sch_multiq.c 8.8 KB
Newer Older
1
// SPDX-License-Identifier: GPL-2.0-only
2 3 4 5 6 7 8
/*
 * Copyright (c) 2008, Intel Corporation.
 *
 * Author: Alexander Duyck <alexander.h.duyck@intel.com>
 */

#include <linux/module.h>
9
#include <linux/slab.h>
10 11 12 13 14 15 16
#include <linux/types.h>
#include <linux/kernel.h>
#include <linux/string.h>
#include <linux/errno.h>
#include <linux/skbuff.h>
#include <net/netlink.h>
#include <net/pkt_sched.h>
17
#include <net/pkt_cls.h>
18 19 20 21 22

struct multiq_sched_data {
	u16 bands;
	u16 max_bands;
	u16 curband;
J
John Fastabend 已提交
23
	struct tcf_proto __rcu *filter_list;
24
	struct tcf_block *block;
25 26 27 28 29 30 31 32 33 34
	struct Qdisc **queues;
};


static struct Qdisc *
multiq_classify(struct sk_buff *skb, struct Qdisc *sch, int *qerr)
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	u32 band;
	struct tcf_result res;
J
John Fastabend 已提交
35
	struct tcf_proto *fl = rcu_dereference_bh(q->filter_list);
36 37 38
	int err;

	*qerr = NET_XMIT_SUCCESS | __NET_XMIT_BYPASS;
39
	err = tcf_classify(skb, fl, &res, false);
40 41 42 43
#ifdef CONFIG_NET_CLS_ACT
	switch (err) {
	case TC_ACT_STOLEN:
	case TC_ACT_QUEUED:
44
	case TC_ACT_TRAP:
45
		*qerr = NET_XMIT_SUCCESS | __NET_XMIT_STOLEN;
46
		/* fall through */
47 48 49 50 51 52 53 54 55 56 57 58 59
	case TC_ACT_SHOT:
		return NULL;
	}
#endif
	band = skb_get_queue_mapping(skb);

	if (band >= q->bands)
		return q->queues[0];

	return q->queues[band];
}

static int
60 61
multiq_enqueue(struct sk_buff *skb, struct Qdisc *sch,
	       struct sk_buff **to_free)
62 63 64 65 66 67 68 69 70
{
	struct Qdisc *qdisc;
	int ret;

	qdisc = multiq_classify(skb, sch, &ret);
#ifdef CONFIG_NET_CLS_ACT
	if (qdisc == NULL) {

		if (ret & __NET_XMIT_BYPASS)
71
			qdisc_qstats_drop(sch);
72
		__qdisc_drop(skb, to_free);
73 74 75 76
		return ret;
	}
#endif

77
	ret = qdisc_enqueue(skb, qdisc, to_free);
78 79 80 81 82
	if (ret == NET_XMIT_SUCCESS) {
		sch->q.qlen++;
		return NET_XMIT_SUCCESS;
	}
	if (net_xmit_drop_count(ret))
83
		qdisc_qstats_drop(sch);
84 85 86 87 88 89 90 91 92 93 94 95 96 97 98 99 100
	return ret;
}

static struct sk_buff *multiq_dequeue(struct Qdisc *sch)
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	struct Qdisc *qdisc;
	struct sk_buff *skb;
	int band;

	for (band = 0; band < q->bands; band++) {
		/* cycle through bands to ensure fairness */
		q->curband++;
		if (q->curband >= q->bands)
			q->curband = 0;

		/* Check that target subqueue is available before
101
		 * pulling an skb to avoid head-of-line blocking.
102
		 */
103 104
		if (!netif_xmit_stopped(
		    netdev_get_tx_queue(qdisc_dev(sch), q->curband))) {
105 106 107
			qdisc = q->queues[q->curband];
			skb = qdisc->dequeue(qdisc);
			if (skb) {
108
				qdisc_bstats_update(sch, skb);
109 110 111 112 113 114 115 116 117
				sch->q.qlen--;
				return skb;
			}
		}
	}
	return NULL;

}

118 119 120 121 122 123 124 125 126 127 128 129 130 131 132
static struct sk_buff *multiq_peek(struct Qdisc *sch)
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	unsigned int curband = q->curband;
	struct Qdisc *qdisc;
	struct sk_buff *skb;
	int band;

	for (band = 0; band < q->bands; band++) {
		/* cycle through bands to ensure fairness */
		curband++;
		if (curband >= q->bands)
			curband = 0;

		/* Check that target subqueue is available before
133
		 * pulling an skb to avoid head-of-line blocking.
134
		 */
135 136
		if (!netif_xmit_stopped(
		    netdev_get_tx_queue(qdisc_dev(sch), curband))) {
137 138 139 140 141 142 143 144 145 146
			qdisc = q->queues[curband];
			skb = qdisc->ops->peek(qdisc);
			if (skb)
				return skb;
		}
	}
	return NULL;

}

147 148 149 150 151 152 153 154 155 156 157 158 159 160 161 162 163 164
static void
multiq_reset(struct Qdisc *sch)
{
	u16 band;
	struct multiq_sched_data *q = qdisc_priv(sch);

	for (band = 0; band < q->bands; band++)
		qdisc_reset(q->queues[band]);
	sch->q.qlen = 0;
	q->curband = 0;
}

static void
multiq_destroy(struct Qdisc *sch)
{
	int band;
	struct multiq_sched_data *q = qdisc_priv(sch);

165
	tcf_block_put(q->block);
166
	for (band = 0; band < q->bands; band++)
167
		qdisc_put(q->queues[band]);
168 169 170 171

	kfree(q->queues);
}

172 173
static int multiq_tune(struct Qdisc *sch, struct nlattr *opt,
		       struct netlink_ext_ack *extack)
174 175 176
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	struct tc_multiq_qopt *qopt;
177 178
	struct Qdisc **removed;
	int i, n_removed = 0;
179 180

	if (!netif_is_multiqueue(qdisc_dev(sch)))
181
		return -EOPNOTSUPP;
182 183 184 185 186 187 188
	if (nla_len(opt) < sizeof(*qopt))
		return -EINVAL;

	qopt = nla_data(opt);

	qopt->bands = qdisc_dev(sch)->real_num_tx_queues;

189 190 191 192 193
	removed = kmalloc(sizeof(*removed) * (q->max_bands - q->bands),
			  GFP_KERNEL);
	if (!removed)
		return -ENOMEM;

194 195 196
	sch_tree_lock(sch);
	q->bands = qopt->bands;
	for (i = q->bands; i < q->max_bands; i++) {
197
		if (q->queues[i] != &noop_qdisc) {
198
			struct Qdisc *child = q->queues[i];
199

200
			q->queues[i] = &noop_qdisc;
201 202
			qdisc_purge_queue(child);
			removed[n_removed++] = child;
203 204 205 206 207
		}
	}

	sch_tree_unlock(sch);

208 209 210 211
	for (i = 0; i < n_removed; i++)
		qdisc_put(removed[i]);
	kfree(removed);

212 213
	for (i = 0; i < q->bands; i++) {
		if (q->queues[i] == &noop_qdisc) {
214
			struct Qdisc *child, *old;
215
			child = qdisc_create_dflt(sch->dev_queue,
216 217
						  &pfifo_qdisc_ops,
						  TC_H_MAKE(sch->handle,
218
							    i + 1), extack);
219 220
			if (child) {
				sch_tree_lock(sch);
221 222
				old = q->queues[i];
				q->queues[i] = child;
223 224
				if (child != &noop_qdisc)
					qdisc_hash_add(child, true);
225

226 227
				if (old != &noop_qdisc)
					qdisc_purge_queue(old);
228
				sch_tree_unlock(sch);
229
				qdisc_put(old);
230 231 232 233 234 235
			}
		}
	}
	return 0;
}

236 237
static int multiq_init(struct Qdisc *sch, struct nlattr *opt,
		       struct netlink_ext_ack *extack)
238 239
{
	struct multiq_sched_data *q = qdisc_priv(sch);
240
	int i, err;
241 242 243

	q->queues = NULL;

244
	if (!opt)
245 246
		return -EINVAL;

247
	err = tcf_block_get(&q->block, &q->filter_list, sch, extack);
248 249 250
	if (err)
		return err;

251 252 253 254 255 256 257 258
	q->max_bands = qdisc_dev(sch)->num_tx_queues;

	q->queues = kcalloc(q->max_bands, sizeof(struct Qdisc *), GFP_KERNEL);
	if (!q->queues)
		return -ENOBUFS;
	for (i = 0; i < q->max_bands; i++)
		q->queues[i] = &noop_qdisc;

259
	return multiq_tune(sch, opt, extack);
260 261 262 263 264 265 266 267 268 269 270
}

static int multiq_dump(struct Qdisc *sch, struct sk_buff *skb)
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	unsigned char *b = skb_tail_pointer(skb);
	struct tc_multiq_qopt opt;

	opt.bands = q->bands;
	opt.max_bands = q->max_bands;

271 272
	if (nla_put(skb, TCA_OPTIONS, sizeof(opt), &opt))
		goto nla_put_failure;
273 274 275 276 277 278 279 280 281

	return skb->len;

nla_put_failure:
	nlmsg_trim(skb, b);
	return -1;
}

static int multiq_graft(struct Qdisc *sch, unsigned long arg, struct Qdisc *new,
282
			struct Qdisc **old, struct netlink_ext_ack *extack)
283 284 285 286 287 288 289
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	unsigned long band = arg - 1;

	if (new == NULL)
		new = &noop_qdisc;

290
	*old = qdisc_replace(sch, new, &q->queues[band]);
291 292 293 294 295 296 297 298 299 300 301 302
	return 0;
}

static struct Qdisc *
multiq_leaf(struct Qdisc *sch, unsigned long arg)
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	unsigned long band = arg - 1;

	return q->queues[band];
}

303
static unsigned long multiq_find(struct Qdisc *sch, u32 classid)
304 305 306 307 308 309 310 311 312 313 314 315
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	unsigned long band = TC_H_MIN(classid);

	if (band - 1 >= q->bands)
		return 0;
	return band;
}

static unsigned long multiq_bind(struct Qdisc *sch, unsigned long parent,
				 u32 classid)
{
316
	return multiq_find(sch, classid);
317 318 319
}


320
static void multiq_unbind(struct Qdisc *q, unsigned long cl)
321 322 323 324 325 326 327 328 329
{
}

static int multiq_dump_class(struct Qdisc *sch, unsigned long cl,
			     struct sk_buff *skb, struct tcmsg *tcm)
{
	struct multiq_sched_data *q = qdisc_priv(sch);

	tcm->tcm_handle |= TC_H_MIN(cl);
E
Eric Dumazet 已提交
330
	tcm->tcm_info = q->queues[cl - 1]->handle;
331 332 333 334 335 336 337 338 339 340
	return 0;
}

static int multiq_dump_class_stats(struct Qdisc *sch, unsigned long cl,
				 struct gnet_dump *d)
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	struct Qdisc *cl_q;

	cl_q = q->queues[cl - 1];
341
	if (gnet_stats_copy_basic(qdisc_root_sleeping_running(sch),
342
				  d, cl_q->cpu_bstats, &cl_q->bstats) < 0 ||
343
	    qdisc_qstats_copy(d, cl_q) < 0)
344 345 346 347 348 349 350 351 352 353 354 355 356 357 358 359 360 361
		return -1;

	return 0;
}

static void multiq_walk(struct Qdisc *sch, struct qdisc_walker *arg)
{
	struct multiq_sched_data *q = qdisc_priv(sch);
	int band;

	if (arg->stop)
		return;

	for (band = 0; band < q->bands; band++) {
		if (arg->count < arg->skip) {
			arg->count++;
			continue;
		}
E
Eric Dumazet 已提交
362
		if (arg->fn(sch, band + 1, arg) < 0) {
363 364 365 366 367 368 369
			arg->stop = 1;
			break;
		}
		arg->count++;
	}
}

370 371
static struct tcf_block *multiq_tcf_block(struct Qdisc *sch, unsigned long cl,
					  struct netlink_ext_ack *extack)
372 373 374 375 376
{
	struct multiq_sched_data *q = qdisc_priv(sch);

	if (cl)
		return NULL;
377
	return q->block;
378 379 380 381 382
}

static const struct Qdisc_class_ops multiq_class_ops = {
	.graft		=	multiq_graft,
	.leaf		=	multiq_leaf,
383
	.find		=	multiq_find,
384
	.walk		=	multiq_walk,
385
	.tcf_block	=	multiq_tcf_block,
386
	.bind_tcf	=	multiq_bind,
387
	.unbind_tcf	=	multiq_unbind,
388 389 390 391 392 393 394 395 396 397 398
	.dump		=	multiq_dump_class,
	.dump_stats	=	multiq_dump_class_stats,
};

static struct Qdisc_ops multiq_qdisc_ops __read_mostly = {
	.next		=	NULL,
	.cl_ops		=	&multiq_class_ops,
	.id		=	"multiq",
	.priv_size	=	sizeof(struct multiq_sched_data),
	.enqueue	=	multiq_enqueue,
	.dequeue	=	multiq_dequeue,
399
	.peek		=	multiq_peek,
400 401 402 403 404 405 406 407 408 409 410 411 412 413 414 415 416 417 418 419 420 421
	.init		=	multiq_init,
	.reset		=	multiq_reset,
	.destroy	=	multiq_destroy,
	.change		=	multiq_tune,
	.dump		=	multiq_dump,
	.owner		=	THIS_MODULE,
};

static int __init multiq_module_init(void)
{
	return register_qdisc(&multiq_qdisc_ops);
}

static void __exit multiq_module_exit(void)
{
	unregister_qdisc(&multiq_qdisc_ops);
}

module_init(multiq_module_init)
module_exit(multiq_module_exit)

MODULE_LICENSE("GPL");