Commit ffb40259 authored by NeilBrown's avatar NeilBrown Committed by Chuck Lever

nfsd: Don't leave work of closing files to a work queue

The work of closing a file can have non-trivial cost.  Doing it in a
separate work queue thread means that cost isn't imposed on the nfsd
threads and an imbalance can be created.  This can result in files being
queued for the work queue more quickly that the work queue can process
them, resulting in unbounded growth of the queue and memory exhaustion.

To avoid this work imbalance that exhausts memory, this patch moves all
closing of files into the nfsd threads.  This means that when the work
imposes a cost, that cost appears where it would be expected - in the
work of the nfsd thread.  A subsequent patch will ensure the final
__fput() is called in the same (nfsd) thread which calls filp_close().

Files opened for NFSv3 are never explicitly closed by the client and are
kept open by the server in the "filecache", which responds to memory
pressure, is garbage collected even when there is no pressure, and
sometimes closes files when there is particular need such as for rename.
These files currently have filp_close() called in a dedicated work
queue, so their __fput() can have no effect on nfsd threads.

This patch discards the work queue and instead has each nfsd thread call
flip_close() on as many as 8 files from the filecache each time it acts
on a client request (or finds there are no pending client requests).  If
there are more to be closed, more threads are woken.  This spreads the
work of __fput() over multiple threads and imposes any cost on those
threads.

The number 8 is somewhat arbitrary.  It needs to be greater than 1 to
ensure that files are closed more quickly than they can be added to the
cache.  It needs to be small enough to limit the per-request delays that
will be imposed on clients when all threads are busy closing files.
Signed-off-by: default avatarNeilBrown <neilb@suse.de>
Reviewed-by: default avatarJeff Layton <jlayton@kernel.org>
Signed-off-by: default avatarChuck Lever <chuck.lever@oracle.com>
parent 561141dd
...@@ -61,13 +61,10 @@ static DEFINE_PER_CPU(unsigned long, nfsd_file_total_age); ...@@ -61,13 +61,10 @@ static DEFINE_PER_CPU(unsigned long, nfsd_file_total_age);
static DEFINE_PER_CPU(unsigned long, nfsd_file_evictions); static DEFINE_PER_CPU(unsigned long, nfsd_file_evictions);
struct nfsd_fcache_disposal { struct nfsd_fcache_disposal {
struct work_struct work;
spinlock_t lock; spinlock_t lock;
struct list_head freeme; struct list_head freeme;
}; };
static struct workqueue_struct *nfsd_filecache_wq __read_mostly;
static struct kmem_cache *nfsd_file_slab; static struct kmem_cache *nfsd_file_slab;
static struct kmem_cache *nfsd_file_mark_slab; static struct kmem_cache *nfsd_file_mark_slab;
static struct list_lru nfsd_file_lru; static struct list_lru nfsd_file_lru;
...@@ -421,7 +418,37 @@ nfsd_file_dispose_list_delayed(struct list_head *dispose) ...@@ -421,7 +418,37 @@ nfsd_file_dispose_list_delayed(struct list_head *dispose)
spin_lock(&l->lock); spin_lock(&l->lock);
list_move_tail(&nf->nf_lru, &l->freeme); list_move_tail(&nf->nf_lru, &l->freeme);
spin_unlock(&l->lock); spin_unlock(&l->lock);
queue_work(nfsd_filecache_wq, &l->work); svc_wake_up(nn->nfsd_serv);
}
}
/**
* nfsd_file_net_dispose - deal with nfsd_files waiting to be disposed.
* @nn: nfsd_net in which to find files to be disposed.
*
* When files held open for nfsv3 are removed from the filecache, whether
* due to memory pressure or garbage collection, they are queued to
* a per-net-ns queue. This function completes the disposal, either
* directly or by waking another nfsd thread to help with the work.
*/
void nfsd_file_net_dispose(struct nfsd_net *nn)
{
struct nfsd_fcache_disposal *l = nn->fcache_disposal;
if (!list_empty(&l->freeme)) {
LIST_HEAD(dispose);
int i;
spin_lock(&l->lock);
for (i = 0; i < 8 && !list_empty(&l->freeme); i++)
list_move(l->freeme.next, &dispose);
spin_unlock(&l->lock);
if (!list_empty(&l->freeme))
/* Wake up another thread to share the work
* *before* doing any actual disposing.
*/
svc_wake_up(nn->nfsd_serv);
nfsd_file_dispose_list(&dispose);
} }
} }
...@@ -634,27 +661,6 @@ nfsd_file_close_inode_sync(struct inode *inode) ...@@ -634,27 +661,6 @@ nfsd_file_close_inode_sync(struct inode *inode)
flush_delayed_fput(); flush_delayed_fput();
} }
/**
* nfsd_file_delayed_close - close unused nfsd_files
* @work: dummy
*
* Scrape the freeme list for this nfsd_net, and then dispose of them
* all.
*/
static void
nfsd_file_delayed_close(struct work_struct *work)
{
LIST_HEAD(head);
struct nfsd_fcache_disposal *l = container_of(work,
struct nfsd_fcache_disposal, work);
spin_lock(&l->lock);
list_splice_init(&l->freeme, &head);
spin_unlock(&l->lock);
nfsd_file_dispose_list(&head);
}
static int static int
nfsd_file_lease_notifier_call(struct notifier_block *nb, unsigned long arg, nfsd_file_lease_notifier_call(struct notifier_block *nb, unsigned long arg,
void *data) void *data)
...@@ -717,10 +723,6 @@ nfsd_file_cache_init(void) ...@@ -717,10 +723,6 @@ nfsd_file_cache_init(void)
return ret; return ret;
ret = -ENOMEM; ret = -ENOMEM;
nfsd_filecache_wq = alloc_workqueue("nfsd_filecache", WQ_UNBOUND, 0);
if (!nfsd_filecache_wq)
goto out;
nfsd_file_slab = kmem_cache_create("nfsd_file", nfsd_file_slab = kmem_cache_create("nfsd_file",
sizeof(struct nfsd_file), 0, 0, NULL); sizeof(struct nfsd_file), 0, 0, NULL);
if (!nfsd_file_slab) { if (!nfsd_file_slab) {
...@@ -735,7 +737,6 @@ nfsd_file_cache_init(void) ...@@ -735,7 +737,6 @@ nfsd_file_cache_init(void)
goto out_err; goto out_err;
} }
ret = list_lru_init(&nfsd_file_lru); ret = list_lru_init(&nfsd_file_lru);
if (ret) { if (ret) {
pr_err("nfsd: failed to init nfsd_file_lru: %d\n", ret); pr_err("nfsd: failed to init nfsd_file_lru: %d\n", ret);
...@@ -785,8 +786,6 @@ nfsd_file_cache_init(void) ...@@ -785,8 +786,6 @@ nfsd_file_cache_init(void)
nfsd_file_slab = NULL; nfsd_file_slab = NULL;
kmem_cache_destroy(nfsd_file_mark_slab); kmem_cache_destroy(nfsd_file_mark_slab);
nfsd_file_mark_slab = NULL; nfsd_file_mark_slab = NULL;
destroy_workqueue(nfsd_filecache_wq);
nfsd_filecache_wq = NULL;
rhltable_destroy(&nfsd_file_rhltable); rhltable_destroy(&nfsd_file_rhltable);
goto out; goto out;
} }
...@@ -832,7 +831,6 @@ nfsd_alloc_fcache_disposal(void) ...@@ -832,7 +831,6 @@ nfsd_alloc_fcache_disposal(void)
l = kmalloc(sizeof(*l), GFP_KERNEL); l = kmalloc(sizeof(*l), GFP_KERNEL);
if (!l) if (!l)
return NULL; return NULL;
INIT_WORK(&l->work, nfsd_file_delayed_close);
spin_lock_init(&l->lock); spin_lock_init(&l->lock);
INIT_LIST_HEAD(&l->freeme); INIT_LIST_HEAD(&l->freeme);
return l; return l;
...@@ -841,7 +839,6 @@ nfsd_alloc_fcache_disposal(void) ...@@ -841,7 +839,6 @@ nfsd_alloc_fcache_disposal(void)
static void static void
nfsd_free_fcache_disposal(struct nfsd_fcache_disposal *l) nfsd_free_fcache_disposal(struct nfsd_fcache_disposal *l)
{ {
cancel_work_sync(&l->work);
nfsd_file_dispose_list(&l->freeme); nfsd_file_dispose_list(&l->freeme);
kfree(l); kfree(l);
} }
...@@ -910,8 +907,6 @@ nfsd_file_cache_shutdown(void) ...@@ -910,8 +907,6 @@ nfsd_file_cache_shutdown(void)
fsnotify_wait_marks_destroyed(); fsnotify_wait_marks_destroyed();
kmem_cache_destroy(nfsd_file_mark_slab); kmem_cache_destroy(nfsd_file_mark_slab);
nfsd_file_mark_slab = NULL; nfsd_file_mark_slab = NULL;
destroy_workqueue(nfsd_filecache_wq);
nfsd_filecache_wq = NULL;
rhltable_destroy(&nfsd_file_rhltable); rhltable_destroy(&nfsd_file_rhltable);
for_each_possible_cpu(i) { for_each_possible_cpu(i) {
......
...@@ -56,6 +56,7 @@ void nfsd_file_cache_shutdown_net(struct net *net); ...@@ -56,6 +56,7 @@ void nfsd_file_cache_shutdown_net(struct net *net);
void nfsd_file_put(struct nfsd_file *nf); void nfsd_file_put(struct nfsd_file *nf);
struct nfsd_file *nfsd_file_get(struct nfsd_file *nf); struct nfsd_file *nfsd_file_get(struct nfsd_file *nf);
void nfsd_file_close_inode_sync(struct inode *inode); void nfsd_file_close_inode_sync(struct inode *inode);
void nfsd_file_net_dispose(struct nfsd_net *nn);
bool nfsd_file_is_cached(struct inode *inode); bool nfsd_file_is_cached(struct inode *inode);
__be32 nfsd_file_acquire_gc(struct svc_rqst *rqstp, struct svc_fh *fhp, __be32 nfsd_file_acquire_gc(struct svc_rqst *rqstp, struct svc_fh *fhp,
unsigned int may_flags, struct nfsd_file **nfp); unsigned int may_flags, struct nfsd_file **nfp);
......
...@@ -941,6 +941,8 @@ nfsd(void *vrqstp) ...@@ -941,6 +941,8 @@ nfsd(void *vrqstp)
rqstp->rq_server->sv_maxconn = nn->max_connections; rqstp->rq_server->sv_maxconn = nn->max_connections;
svc_recv(rqstp); svc_recv(rqstp);
nfsd_file_net_dispose(nn);
} }
atomic_dec(&nfsdstats.th_cnt); atomic_dec(&nfsdstats.th_cnt);
......
Markdown is supported
0%
or
You are about to add 0 people to the discussion. Proceed with caution.
Finish editing this message first!
Please register or to comment