Linux-EROFS Archive on lore.kernel.org
 help / color / mirror / Atom feed
From: David Howells <dhowells@redhat.com>
To: Christian Brauner <christian@brauner.io>,
	Matthew Wilcox <willy@infradead.org>,
	Christoph Hellwig <hch@infradead.org>
Cc: David Howells <dhowells@redhat.com>,
	Paulo Alcantara <pc@manguebit.org>, Jens Axboe <axboe@kernel.dk>,
	Leon Romanovsky <leon@kernel.org>,
	Steve French <sfrench@samba.org>,
	ChenXiaoSong <chenxiaosong@chenxiaosong.com>,
	Marc Dionne <marc.dionne@auristor.com>,
	Stefan Metzmacher <metze@samba.org>,
	Eric Van Hensbergen <ericvh@kernel.org>,
	Dominique Martinet <asmadeus@codewreck.org>,
	Ilya Dryomov <idryomov@gmail.com>,
	netfs@lists.linux.dev, linux-afs@lists.infradead.org,
	linux-cifs@vger.kernel.org, linux-nfs@vger.kernel.org,
	ceph-devel@vger.kernel.org, v9fs@lists.linux.dev,
	linux-erofs@lists.ozlabs.org, linux-fsdevel@vger.kernel.org,
	linux-kernel@vger.kernel.org
Subject: [PATCH v8 07/25] netfs: Make mempool available for bvecq
Date: Tue,  4 Aug 2026 11:02:02 +0100	[thread overview]
Message-ID: <20260804100224.2748935-8-dhowells@redhat.com> (raw)
In-Reply-To: <20260804100224.2748935-1-dhowells@redhat.com>

Make a mempool available for allocating bvecq structs.  Use it
automatically if anything other than GFP_KERNEL (with GFP_ZONEMASK flags
masked off) is specified.  Reading from a file can use GFP_KERNEL as the
failure mode is straightforward and the same for DIO reads and writes.
When it comes to writeback, however, the writeback_iter() API does not
permit temporary failure, including ENOMEM, if WB_SYNC_ALL is set and the
caller must process all folios to completion.  (I'm not sure if EINTR
constitutes an acceptable failure).

Signed-off-by: David Howells <dhowells@redhat.com>
cc: Paulo Alcantara <pc@manguebit.org>
cc: Matthew Wilcox <willy@infradead.org>
cc: Christoph Hellwig <hch@infradead.org>
cc: linux-cifs@vger.kernel.org
cc: netfs@lists.linux.dev
cc: linux-fsdevel@vger.kernel.org
---
 fs/netfs/bvecq.c      | 27 ++++++++++++++++++++++-----
 fs/netfs/internal.h   |  1 +
 fs/netfs/main.c       |  7 +++++++
 include/linux/bvecq.h |  3 +++
 4 files changed, 33 insertions(+), 5 deletions(-)

diff --git a/fs/netfs/bvecq.c b/fs/netfs/bvecq.c
index b8dd5250ec05..9be1d07a158e 100644
--- a/fs/netfs/bvecq.c
+++ b/fs/netfs/bvecq.c
@@ -45,8 +45,9 @@ EXPORT_SYMBOL(bvecq_dump);
  *
  * Allocate a single bvecq node and initialise the header.  A number of inline
  * slots are also allocated, rounded up to fit after the header in a power-of-2
- * slab object of up to 512 bytes (up to 29 slots on a 64-bit cpu).  The slot
- * array is not initialised.
+ * slab object of up to 512 bytes (up to 29 slots on a 64-bit cpu).  The caller
+ * should be aware that the number of slots allocated may be more or less than
+ * the number requested.  The slot array is not initialised.
  *
  * Return: The node pointer or NULL on allocation failure.
  */
@@ -57,14 +58,27 @@ struct bvecq *bvecq_alloc_one(size_t nr_slots, gfp_t gfp)
 	const size_t max_slots = (max_size - sizeof(*bq)) / sizeof(bq->__bv[0]);
 	size_t part = umin(nr_slots, max_slots);
 	size_t size = roundup_pow_of_two(struct_size(bq, __bv, part));
-
-	bq = kmalloc(size, gfp & ~GFP_ZONEMASK);
+	bool from_pool = false;
+
+	gfp &= ~GFP_ZONEMASK;
+	if (size != BVECQ_STD_SIZE) {
+		bq = kmalloc(size, gfp);
+	} else {
+		bq = netfs_bvecq_pool.alloc(gfp, netfs_bvecq_pool.pool_data);
+		from_pool = true;
+	}
+	if (!bq && gfp != GFP_KERNEL) {
+		bq = mempool_alloc(&netfs_bvecq_pool, gfp);
+		from_pool = true;
+		size = BVECQ_STD_SIZE;
+	}
 	if (bq) {
 		*bq = (struct bvecq) {
 			.ref		= REFCOUNT_INIT(1),
 			.bv		= bq->__bv,
 			.inline_bv	= true,
 			.max_slots	= (size - sizeof(*bq)) / sizeof(bq->__bv[0]),
+			.from_pool	= from_pool,
 		};
 		netfs_stat(&netfs_n_bvecq);
 	}
@@ -242,7 +256,10 @@ void bvecq_put(struct bvecq *bq)
 			bvecq_free_slot(bq, slot);
 		next = bq->next;
 		netfs_stat_d(&netfs_n_bvecq);
-		kfree(bq);
+		if (bq->from_pool)
+			mempool_free(bq, &netfs_bvecq_pool);
+		else
+			kfree(bq);
 	}
 }
 EXPORT_SYMBOL(bvecq_put);
diff --git a/fs/netfs/internal.h b/fs/netfs/internal.h
index 46fb92601774..c68eea2ecc60 100644
--- a/fs/netfs/internal.h
+++ b/fs/netfs/internal.h
@@ -43,6 +43,7 @@ extern struct list_head netfs_io_requests;
 extern spinlock_t netfs_proc_lock;
 extern mempool_t netfs_request_pool;
 extern mempool_t netfs_subrequest_pool;
+extern mempool_t netfs_bvecq_pool;
 extern mempool_t netfs_folioq_pool;
 
 #ifdef CONFIG_PROC_FS
diff --git a/fs/netfs/main.c b/fs/netfs/main.c
index 927badf3989d..9f72e5054aff 100644
--- a/fs/netfs/main.c
+++ b/fs/netfs/main.c
@@ -28,6 +28,7 @@ static struct kmem_cache *netfs_request_slab;
 static struct kmem_cache *netfs_subrequest_slab;
 mempool_t netfs_request_pool;
 mempool_t netfs_subrequest_pool;
+mempool_t netfs_bvecq_pool;
 mempool_t netfs_folioq_pool;
 
 #ifdef CONFIG_PROC_FS
@@ -112,6 +113,9 @@ static int __init netfs_init(void)
 	if (mempool_init_kmalloc_pool(&netfs_folioq_pool, 100, sizeof(struct folio_queue)) < 0)
 		goto error_folioq_pool;
 
+	if (mempool_init_kmalloc_pool(&netfs_bvecq_pool, 100, BVECQ_STD_SIZE) < 0)
+		goto error_bvecq_pool;
+
 	netfs_request_slab = kmem_cache_create("netfs_request",
 					       sizeof(struct netfs_io_request), 0,
 					       SLAB_HWCACHE_ALIGN | SLAB_ACCOUNT,
@@ -164,6 +168,8 @@ static int __init netfs_init(void)
 error_reqpool:
 	kmem_cache_destroy(netfs_request_slab);
 error_req:
+	mempool_exit(&netfs_bvecq_pool);
+error_bvecq_pool:
 	mempool_exit(&netfs_folioq_pool);
 error_folioq_pool:
 	return ret;
@@ -178,6 +184,7 @@ static void __exit netfs_exit(void)
 	kmem_cache_destroy(netfs_subrequest_slab);
 	mempool_exit(&netfs_request_pool);
 	kmem_cache_destroy(netfs_request_slab);
+	mempool_exit(&netfs_bvecq_pool);
 	mempool_exit(&netfs_folioq_pool);
 }
 module_exit(netfs_exit);
diff --git a/include/linux/bvecq.h b/include/linux/bvecq.h
index fa46be520649..d1476393942e 100644
--- a/include/linux/bvecq.h
+++ b/include/linux/bvecq.h
@@ -49,15 +49,18 @@ struct bvecq {
 	enum bvecq_mem	mem_type:3;	/* What sort of memory and how to free it */
 	bool		inline_bv:1;	/* T if __bv[] is being used */
 	bool		discontig:1;	/* T if not contiguous with previous bvecq */
+	bool		from_pool:1;	/* T if bvecq from mempool */
 	struct bio_vec	*bv;		/* Pointer to array of page fragments */
 	struct bio_vec	__bv[];		/* Default array (if ->inline_bv) */
 };
 
 #if BITS_PER_LONG == 64
 /* Number of slots in __bv[] for a bvecq in a 512-byte kmalloc block. */
+#define BVECQ_STD_SIZE		512
 #define BVECQ_STD_SLOTS		29	/* 2 words/slot; 32 slots; bvecq is 6 words (3 slots) */
 #elif  BITS_PER_LONG == 32
 /* Number of slots in __bv[] for a bvecq in a 256-byte kmalloc block. */
+#define BVECQ_STD_SIZE		256
 #define BVECQ_STD_SLOTS		18	/* 3 words/slot; 21 slots; bvecq is 9 words (3 slots) */
 #else
 #error BVECQ_STD_SLOTS undetermined



  parent reply	other threads:[~2026-08-04 10:03 UTC|newest]

Thread overview: 26+ messages / expand[flat|nested]  mbox.gz  Atom feed  top
2026-08-04 10:01 [PATCH v8 00/25] netfs: Keep track of folios in a segmented bio_vec[] chain David Howells
2026-08-04 10:01 ` [PATCH v8 01/25] mm: Make readahead store folio count in readahead_control David Howells
2026-08-04 10:01 ` [PATCH v8 02/25] netfs: Bulk load the readahead-provided folios up front David Howells
2026-08-04 10:01 ` [PATCH v8 03/25] Add a function to kmap one page of a multipage bio_vec David Howells
2026-08-04 10:01 ` [PATCH v8 04/25] iov_iter: Make iov_iter_get_pages*() wrap iov_iter_extract_pages() David Howells
2026-08-04 10:02 ` [PATCH v8 05/25] iov_iter: Add a segmented queue of bio_vec[] David Howells
2026-08-04 10:02 ` [PATCH v8 06/25] netfs: Add some tools for managing bvecq chains David Howells
2026-08-04 10:02 ` David Howells [this message]
2026-08-04 10:02 ` [PATCH v8 08/25] netfs: Add a function to extract from an iter into a bvecq David Howells
2026-08-04 10:02 ` [PATCH v8 09/25] afs: Use a bvecq to hold dir content rather than folioq David Howells
2026-08-04 10:02 ` [PATCH v8 10/25] cifs: Use a bvecq for buffering instead of a folioq David Howells
2026-08-04 10:02 ` [PATCH v8 11/25] smbdirect: Support ITER_BVECQ in smbdirect_map_sges_from_iter() David Howells
2026-08-04 10:02 ` [PATCH v8 12/25] netfs: Remove the writethrough code David Howells
2026-08-04 10:02 ` [PATCH v8 13/25] cachefiles,netfs: sunset ondemand mode David Howells
2026-08-04 10:02 ` [PATCH v8 14/25] cachefiles: Don't rely on backing fs storage map for most use cases David Howells
2026-08-04 10:02 ` [PATCH v8 15/25] netfs: Add the cache object ID to netfs_read/write tracepoints David Howells
2026-08-04 10:02 ` [PATCH v8 16/25] netfs: Switch to using bvecq rather than folio_queue and rolling_buffer David Howells
2026-08-04 10:02 ` [PATCH v8 17/25] smbdirect: Remove support for ITER_FOLIOQ from smbdirect_map_sges_from_iter() David Howells
2026-08-04 10:02 ` [PATCH v8 18/25] netfs: Remove netfs_alloc/free_folioq_buffer() David Howells
2026-08-04 10:02 ` [PATCH v8 19/25] netfs: Remove netfs_extract_user_iter() David Howells
2026-08-04 10:02 ` [PATCH v8 20/25] iov_iter: Remove ITER_FOLIOQ David Howells
2026-08-04 10:02 ` [PATCH v8 21/25] netfs: Remove folio_queue and rolling_buffer David Howells
2026-08-04 10:02 ` [PATCH v8 22/25] netfs: Check for too much data being read David Howells
2026-08-04 10:02 ` [PATCH v8 23/25] netfs: Limit the minimum trigger for progress reporting David Howells
2026-08-04 10:02 ` [PATCH v8 24/25] netfs: Combine prepare and issue ops and grab the buffers on request David Howells
2026-08-04 10:02 ` [PATCH v8 25/25] cachefiles: Preset the state xattr when creating a new file David Howells

Reply instructions:

You may reply publicly to this message via plain-text email
using any one of the following methods:

* Save the following mbox file, import it into your mail client,
  and reply-to-all from there: mbox

  Avoid top-posting and favor interleaved quoting:
  https://en.wikipedia.org/wiki/Posting_style#Interleaved_style

* Reply using the --to, --cc, and --in-reply-to
  switches of git-send-email(1):

  git send-email \
    --in-reply-to=20260804100224.2748935-8-dhowells@redhat.com \
    --to=dhowells@redhat.com \
    --cc=asmadeus@codewreck.org \
    --cc=axboe@kernel.dk \
    --cc=ceph-devel@vger.kernel.org \
    --cc=chenxiaosong@chenxiaosong.com \
    --cc=christian@brauner.io \
    --cc=ericvh@kernel.org \
    --cc=hch@infradead.org \
    --cc=idryomov@gmail.com \
    --cc=leon@kernel.org \
    --cc=linux-afs@lists.infradead.org \
    --cc=linux-cifs@vger.kernel.org \
    --cc=linux-erofs@lists.ozlabs.org \
    --cc=linux-fsdevel@vger.kernel.org \
    --cc=linux-kernel@vger.kernel.org \
    --cc=linux-nfs@vger.kernel.org \
    --cc=marc.dionne@auristor.com \
    --cc=metze@samba.org \
    --cc=netfs@lists.linux.dev \
    --cc=pc@manguebit.org \
    --cc=sfrench@samba.org \
    --cc=v9fs@lists.linux.dev \
    --cc=willy@infradead.org \
    /path/to/YOUR_REPLY

  https://kernel.org/pub/software/scm/git/docs/git-send-email.html

* If your mail client supports setting the In-Reply-To header
  via mailto: links, try the mailto: link
Be sure your reply has a Subject: header at the top and a blank line before the message body.
This is a public inbox, see mirroring instructions
for how to clone and mirror all data and code used for this inbox