bcache: Move sector allocator to alloc.c

Just reorganizing things a bit. Signed-off-by: Kent Overstreet <kmo@daterainc.com>
author: Kent Overstreet <kmo@daterainc.com> 2013-07-24 21:11:11 -0400
committer: Kent Overstreet <kmo@daterainc.com> 2013-11-11 00:56:32 -0500
commit: 2599b53b7b0ea6103d1661dca74d35480cb8fa1f (patch)
tree: 0b277532e73924b68c3da43a999aa5df33824787 /drivers/md/bcache/request.c
parent: 220bb38c21b83e2f7b842f33220bf727093eca89 (diff)
1 files changed, 4 insertions, 182 deletions
diff --git a/drivers/md/bcache/request.c b/drivers/md/bcache/request.c
index 05c7c216f65e..cf7850a7592c 100644
--- a/drivers/md/bcache/request.c
+++ b/drivers/md/bcache/request.c
@@ -255,186 +255,6 @@ static void bch_data_insert_keys(struct closure *cl)
        closure_return(cl);
 }
-struct open_bucket {
-        struct list_head        list;
-        struct task_struct      *last;
-        unsigned                sectors_free;
-        BKEY_PADDED(key);
-};
-void bch_open_buckets_free(struct cache_set *c)
-{
-        struct open_bucket *b;
-        while (!list_empty(&c->data_buckets)) {
-                b = list_first_entry(&c->data_buckets,
-                                     struct open_bucket, list);
-                list_del(&b->list);
-                kfree(b);
-        }
-}
-int bch_open_buckets_alloc(struct cache_set *c)
-{
-        int i;
-        spin_lock_init(&c->data_bucket_lock);
-        for (i = 0; i < 6; i++) {
-                struct open_bucket *b = kzalloc(sizeof(*b), GFP_KERNEL);
-                if (!b)
-                        return -ENOMEM;
-                list_add(&b->list, &c->data_buckets);
-        }
-        return 0;
-}
-/*
- * We keep multiple buckets open for writes, and try to segregate different
- * write streams for better cache utilization: first we look for a bucket where
- * the last write to it was sequential with the current write, and failing that
- * we look for a bucket that was last used by the same task.
- *
- * The ideas is if you've got multiple tasks pulling data into the cache at the
- * same time, you'll get better cache utilization if you try to segregate their
- * data and preserve locality.
- *
- * For example, say you've starting Firefox at the same time you're copying a
- * bunch of files. Firefox will likely end up being fairly hot and stay in the
- * cache awhile, but the data you copied might not be; if you wrote all that
- * data to the same buckets it'd get invalidated at the same time.
- *
- * Both of those tasks will be doing fairly random IO so we can't rely on
- * detecting sequential IO to segregate their data, but going off of the task
- * should be a sane heuristic.
- */
-static struct open_bucket *pick_data_bucket(struct cache_set *c,
-                                            const struct bkey *search,
-                                            struct task_struct *task,
-                                            struct bkey *alloc)
-{
-        struct open_bucket *ret, *ret_task = NULL;
-        list_for_each_entry_reverse(ret, &c->data_buckets, list)
-                if (!bkey_cmp(&ret->key, search))
-                        goto found;
-                else if (ret->last == task)
-                        ret_task = ret;
-        ret = ret_task ?: list_first_entry(&c->data_buckets,
-                                           struct open_bucket, list);
-found:
-        if (!ret->sectors_free && KEY_PTRS(alloc)) {
-                ret->sectors_free = c->sb.bucket_size;
-                bkey_copy(&ret->key, alloc);
-                bkey_init(alloc);
-        }
-        if (!ret->sectors_free)
-                ret = NULL;
-        return ret;
-}
-/*
- * Allocates some space in the cache to write to, and k to point to the newly
- * allocated space, and updates KEY_SIZE(k) and KEY_OFFSET(k) (to point to the
- * end of the newly allocated space).
- *
- * May allocate fewer sectors than @sectors, KEY_SIZE(k) indicates how many
- * sectors were actually allocated.
- *
- * If s->writeback is true, will not fail.
- */
-static bool bch_alloc_sectors(struct data_insert_op *op,
-                              struct bkey *k, unsigned sectors)
-{
-        struct cache_set *c = op->c;
-        struct open_bucket *b;
-        BKEY_PADDED(key) alloc;
-        unsigned i;
-        /*
-         * We might have to allocate a new bucket, which we can't do with a
-         * spinlock held. So if we have to allocate, we drop the lock, allocate
-         * and then retry. KEY_PTRS() indicates whether alloc points to
-         * allocated bucket(s).
-         */
-        bkey_init(&alloc.key);
-        spin_lock(&c->data_bucket_lock);
-        while (!(b = pick_data_bucket(c, k, op->task, &alloc.key))) {
-                unsigned watermark = op->write_prio
-                        ? WATERMARK_MOVINGGC
-                        : WATERMARK_NONE;
-                spin_unlock(&c->data_bucket_lock);
-                if (bch_bucket_alloc_set(c, watermark, &alloc.key,
-                                         1, op->writeback))
-                        return false;
-                spin_lock(&c->data_bucket_lock);
-        }
-        /*
-         * If we had to allocate, we might race and not need to allocate the
-         * second time we call find_data_bucket(). If we allocated a bucket but
-         * didn't use it, drop the refcount bch_bucket_alloc_set() took:
-         */
-        if (KEY_PTRS(&alloc.key))
-                __bkey_put(c, &alloc.key);
-        for (i = 0; i < KEY_PTRS(&b->key); i++)
-                EBUG_ON(ptr_stale(c, &b->key, i));
-        /* Set up the pointer to the space we're allocating: */
-        for (i = 0; i < KEY_PTRS(&b->key); i++)
-                k->ptr[i] = b->key.ptr[i];
-        sectors = min(sectors, b->sectors_free);
-        SET_KEY_OFFSET(k, KEY_OFFSET(k) + sectors);
-        SET_KEY_SIZE(k, sectors);
-        SET_KEY_PTRS(k, KEY_PTRS(&b->key));
-        /*
-         * Move b to the end of the lru, and keep track of what this bucket was
-         * last used for:
-         */
-        list_move_tail(&b->list, &c->data_buckets);
-        bkey_copy_key(&b->key, k);
-        b->last = op->task;
-        b->sectors_free -= sectors;
-        for (i = 0; i < KEY_PTRS(&b->key); i++) {
-                SET_PTR_OFFSET(&b->key, i, PTR_OFFSET(&b->key, i) + sectors);
-                atomic_long_add(sectors,
-                                &PTR_CACHE(c, &b->key, i)->sectors_written);
-        }
-        if (b->sectors_free < c->sb.block_size)
-                b->sectors_free = 0;
-        /*
-         * k takes refcounts on the buckets it points to until it's inserted
-         * into the btree, but if we're done with this bucket we just transfer
-         * get_data_bucket()'s refcount.
-         */
-        if (b->sectors_free)
-                for (i = 0; i < KEY_PTRS(&b->key); i++)
-                        atomic_inc(&PTR_BUCKET(c, &b->key, i)->pin);
-        spin_unlock(&c->data_bucket_lock);
-        return true;
-}
 static void bch_data_invalidate(struct closure *cl)
 {
        struct data_insert_op *op = container_of(cl, struct data_insert_op, cl);
@@ -545,7 +365,9 @@ static void bch_data_insert_start(struct closure *cl)
                SET_KEY_INODE(k, op->inode);
                SET_KEY_OFFSET(k, bio->bi_sector);
-                if (!bch_alloc_sectors(op, k, bio_sectors(bio)))
+                if (!bch_alloc_sectors(op->c, k, bio_sectors(bio),
+                                       op->write_point, op->write_prio,
+                                       op->writeback))
                        goto err;
                n = bch_bio_split(bio, KEY_SIZE(k), GFP_NOIO, split);
@@ -968,7 +790,7 @@ static struct search *search_alloc(struct bio *bio, struct bcache_device *d)
        s->iop.c                = d->c;
        s->d                    = d;
        s->op.lock              = -1;
-        s->iop.task             = current;
+        s->iop.write_point      = hash_long((unsigned long) current, 16);
        s->orig_bio             = bio;
        s->write                = (bio->bi_rw & REQ_WRITE) != 0;
        s->iop.flush_journal    = (bio->bi_rw & (REQ_FLUSH|REQ_FUA)) != 0;
author	Kent Overstreet <kmo@daterainc.com>	2013-07-24 21:11:11 -0400
committer	Kent Overstreet <kmo@daterainc.com>	2013-11-11 00:56:32 -0500
commit	2599b53b7b0ea6103d1661dca74d35480cb8fa1f (patch)
tree	0b277532e73924b68c3da43a999aa5df33824787 /drivers/md/bcache/request.c
parent	220bb38c21b83e2f7b842f33220bf727093eca89 (diff)

diff --git a/drivers/md/bcache/request.c b/drivers/md/bcache/request.c index 05c7c216f65e..cf7850a7592c 100644 --- a/drivers/md/bcache/request.c +++ b/drivers/md/bcache/request.c
@@ -255,186 +255,6 @@ static void bch_data_insert_keys(struct closure *cl)
255	closure_return(cl);	255	closure_return(cl);
256	}	256	}
257		257
258	struct open_bucket {
259	struct list_head list;
260	struct task_struct *last;
261	unsigned sectors_free;
262	BKEY_PADDED(key);
263	};
264
265	void bch_open_buckets_free(struct cache_set *c)
266	{
267	struct open_bucket *b;
268
269	while (!list_empty(&c->data_buckets)) {
270	b = list_first_entry(&c->data_buckets,
271	struct open_bucket, list);
272	list_del(&b->list);
273	kfree(b);
274	}
275	}
276
277	int bch_open_buckets_alloc(struct cache_set *c)
278	{
279	int i;
280
281	spin_lock_init(&c->data_bucket_lock);
282
283	for (i = 0; i < 6; i++) {
284	struct open_bucket b = kzalloc(sizeof(b), GFP_KERNEL);
285	if (!b)
286	return -ENOMEM;
287
288	list_add(&b->list, &c->data_buckets);
289	}
290
291	return 0;
292	}
293
294	/*
295	* We keep multiple buckets open for writes, and try to segregate different
296	* write streams for better cache utilization: first we look for a bucket where
297	* the last write to it was sequential with the current write, and failing that
298	* we look for a bucket that was last used by the same task.
299	*
300	* The ideas is if you've got multiple tasks pulling data into the cache at the
301	* same time, you'll get better cache utilization if you try to segregate their
302	* data and preserve locality.
303	*
304	* For example, say you've starting Firefox at the same time you're copying a
305	* bunch of files. Firefox will likely end up being fairly hot and stay in the
306	* cache awhile, but the data you copied might not be; if you wrote all that
307	* data to the same buckets it'd get invalidated at the same time.
308	*
309	* Both of those tasks will be doing fairly random IO so we can't rely on
310	* detecting sequential IO to segregate their data, but going off of the task
311	* should be a sane heuristic.
312	*/
313	static struct open_bucket pick_data_bucket(struct cache_set c,
314	const struct bkey *search,
315	struct task_struct *task,
316	struct bkey *alloc)
317	{
318	struct open_bucket ret, ret_task = NULL;
319
320	list_for_each_entry_reverse(ret, &c->data_buckets, list)
321	if (!bkey_cmp(&ret->key, search))
322	goto found;
323	else if (ret->last == task)
324	ret_task = ret;
325
326	ret = ret_task ?: list_first_entry(&c->data_buckets,
327	struct open_bucket, list);
328	found:
329	if (!ret->sectors_free && KEY_PTRS(alloc)) {
330	ret->sectors_free = c->sb.bucket_size;
331	bkey_copy(&ret->key, alloc);
332	bkey_init(alloc);
333	}
334
335	if (!ret->sectors_free)
336	ret = NULL;
337
338	return ret;
339	}
340
341	/*
342	* Allocates some space in the cache to write to, and k to point to the newly
343	* allocated space, and updates KEY_SIZE(k) and KEY_OFFSET(k) (to point to the
344	* end of the newly allocated space).
345	*
346	* May allocate fewer sectors than @sectors, KEY_SIZE(k) indicates how many
347	* sectors were actually allocated.
348	*
349	* If s->writeback is true, will not fail.
350	*/
351	static bool bch_alloc_sectors(struct data_insert_op *op,
352	struct bkey *k, unsigned sectors)
353	{
354	struct cache_set *c = op->c;
355	struct open_bucket *b;
356	BKEY_PADDED(key) alloc;
357	unsigned i;
358
359	/*
360	* We might have to allocate a new bucket, which we can't do with a
361	* spinlock held. So if we have to allocate, we drop the lock, allocate
362	* and then retry. KEY_PTRS() indicates whether alloc points to
363	* allocated bucket(s).
364	*/
365
366	bkey_init(&alloc.key);
367	spin_lock(&c->data_bucket_lock);
368
369	while (!(b = pick_data_bucket(c, k, op->task, &alloc.key))) {
370	unsigned watermark = op->write_prio
371	? WATERMARK_MOVINGGC
372	: WATERMARK_NONE;
373
374	spin_unlock(&c->data_bucket_lock);
375
376	if (bch_bucket_alloc_set(c, watermark, &alloc.key,
377	1, op->writeback))
378	return false;
379
380	spin_lock(&c->data_bucket_lock);
381	}
382
383	/*
384	* If we had to allocate, we might race and not need to allocate the
385	* second time we call find_data_bucket(). If we allocated a bucket but
386	* didn't use it, drop the refcount bch_bucket_alloc_set() took:
387	*/
388	if (KEY_PTRS(&alloc.key))
389	__bkey_put(c, &alloc.key);
390
391	for (i = 0; i < KEY_PTRS(&b->key); i++)
392	EBUG_ON(ptr_stale(c, &b->key, i));
393
394	/* Set up the pointer to the space we're allocating: */
395
396	for (i = 0; i < KEY_PTRS(&b->key); i++)
397	k->ptr[i] = b->key.ptr[i];
398
399	sectors = min(sectors, b->sectors_free);
400
401	SET_KEY_OFFSET(k, KEY_OFFSET(k) + sectors);
402	SET_KEY_SIZE(k, sectors);
403	SET_KEY_PTRS(k, KEY_PTRS(&b->key));
404
405	/*
406	* Move b to the end of the lru, and keep track of what this bucket was
407	* last used for:
408	*/
409	list_move_tail(&b->list, &c->data_buckets);
410	bkey_copy_key(&b->key, k);
411	b->last = op->task;
412
413	b->sectors_free -= sectors;
414
415	for (i = 0; i < KEY_PTRS(&b->key); i++) {
416	SET_PTR_OFFSET(&b->key, i, PTR_OFFSET(&b->key, i) + sectors);
417
418	atomic_long_add(sectors,
419	&PTR_CACHE(c, &b->key, i)->sectors_written);
420	}
421
422	if (b->sectors_free < c->sb.block_size)
423	b->sectors_free = 0;
424
425	/*
426	* k takes refcounts on the buckets it points to until it's inserted
427	* into the btree, but if we're done with this bucket we just transfer
428	* get_data_bucket()'s refcount.
429	*/
430	if (b->sectors_free)
431	for (i = 0; i < KEY_PTRS(&b->key); i++)
432	atomic_inc(&PTR_BUCKET(c, &b->key, i)->pin);
433
434	spin_unlock(&c->data_bucket_lock);
435	return true;
436	}
437
438	static void bch_data_invalidate(struct closure *cl)	258	static void bch_data_invalidate(struct closure *cl)
439	{	259	{
440	struct data_insert_op *op = container_of(cl, struct data_insert_op, cl);	260	struct data_insert_op *op = container_of(cl, struct data_insert_op, cl);
@@ -545,7 +365,9 @@ static void bch_data_insert_start(struct closure *cl)
545	SET_KEY_INODE(k, op->inode);	365	SET_KEY_INODE(k, op->inode);
546	SET_KEY_OFFSET(k, bio->bi_sector);	366	SET_KEY_OFFSET(k, bio->bi_sector);
547		367
548	if (!bch_alloc_sectors(op, k, bio_sectors(bio)))	368	if (!bch_alloc_sectors(op->c, k, bio_sectors(bio),
		369	op->write_point, op->write_prio,
		370	op->writeback))
549	goto err;	371	goto err;
550		372
551	n = bch_bio_split(bio, KEY_SIZE(k), GFP_NOIO, split);	373	n = bch_bio_split(bio, KEY_SIZE(k), GFP_NOIO, split);
@@ -968,7 +790,7 @@ static struct search search_alloc(struct bio bio, struct bcache_device *d)
968	s->iop.c = d->c;	790	s->iop.c = d->c;
969	s->d = d;	791	s->d = d;
970	s->op.lock = -1;	792	s->op.lock = -1;
971	s->iop.task = current;	793	s->iop.write_point = hash_long((unsigned long) current, 16);
972	s->orig_bio = bio;	794	s->orig_bio = bio;
973	s->write = (bio->bi_rw & REQ_WRITE) != 0;	795	s->write = (bio->bi_rw & REQ_WRITE) != 0;
974	s->iop.flush_journal = (bio->bi_rw & (REQ_FLUSH\|REQ_FUA)) != 0;	796	s->iop.flush_journal = (bio->bi_rw & (REQ_FLUSH\|REQ_FUA)) != 0;