ipcomp6.c 11.0 KB
Newer Older
L
Linus Torvalds 已提交
1 2 3 4 5 6 7 8 9 10 11
/*
 * IP Payload Compression Protocol (IPComp) for IPv6 - RFC3173
 *
 * Copyright (C)2003 USAGI/WIDE Project
 *
 * Author	Mitsuru KANDA  <mk@linux-ipv6.org>
 *
 * This program is free software; you can redistribute it and/or modify
 * it under the terms of the GNU General Public License as published by
 * the Free Software Foundation; either version 2 of the License, or
 * (at your option) any later version.
12
 *
L
Linus Torvalds 已提交
13 14 15 16
 * This program is distributed in the hope that it will be useful,
 * but WITHOUT ANY WARRANTY; without even the implied warranty of
 * MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE.  See the
 * GNU General Public License for more details.
17
 *
L
Linus Torvalds 已提交
18 19 20 21
 * You should have received a copy of the GNU General Public License
 * along with this program; if not, write to the Free Software
 * Foundation, Inc., 59 Temple Place, Suite 330, Boston, MA  02111-1307  USA
 */
22
/*
L
Linus Torvalds 已提交
23 24 25
 * [Memo]
 *
 * Outbound:
26 27
 *  The compression of IP datagram MUST be done before AH/ESP processing,
 *  fragmentation, and the addition of Hop-by-Hop/Routing header.
L
Linus Torvalds 已提交
28 29
 *
 * Inbound:
30
 *  The decompression of IP datagram MUST be done after the reassembly,
L
Linus Torvalds 已提交
31 32 33 34 35 36 37 38 39 40 41 42 43 44 45 46 47 48
 *  AH/ESP processing.
 */
#include <linux/module.h>
#include <net/ip.h>
#include <net/xfrm.h>
#include <net/ipcomp.h>
#include <asm/scatterlist.h>
#include <asm/semaphore.h>
#include <linux/crypto.h>
#include <linux/pfkeyv2.h>
#include <linux/random.h>
#include <linux/percpu.h>
#include <linux/smp.h>
#include <linux/list.h>
#include <linux/vmalloc.h>
#include <linux/rtnetlink.h>
#include <net/icmp.h>
#include <net/ipv6.h>
49
#include <net/protocol.h>
L
Linus Torvalds 已提交
50 51
#include <linux/ipv6.h>
#include <linux/icmpv6.h>
A
Arjan van de Ven 已提交
52
#include <linux/mutex.h>
L
Linus Torvalds 已提交
53 54 55

struct ipcomp6_tfms {
	struct list_head list;
56
	struct crypto_comp **tfms;
L
Linus Torvalds 已提交
57 58 59
	int users;
};

A
Arjan van de Ven 已提交
60
static DEFINE_MUTEX(ipcomp6_resource_mutex);
L
Linus Torvalds 已提交
61 62 63 64
static void **ipcomp6_scratches;
static int ipcomp6_scratch_users;
static LIST_HEAD(ipcomp6_tfms_list);

65
static int ipcomp6_input(struct xfrm_state *x, struct sk_buff *skb)
L
Linus Torvalds 已提交
66
{
H
Herbert Xu 已提交
67
	int err = -ENOMEM;
L
Linus Torvalds 已提交
68
	struct ipv6hdr *iph;
69
	struct ipv6_comp_hdr *ipch;
L
Linus Torvalds 已提交
70 71 72
	int plen, dlen;
	struct ipcomp_data *ipcd = x->data;
	u8 *start, *scratch;
73
	struct crypto_comp *tfm;
L
Linus Torvalds 已提交
74 75
	int cpu;

H
Herbert Xu 已提交
76
	if (skb_linearize_cow(skb))
L
Linus Torvalds 已提交
77 78 79 80 81
		goto out;

	skb->ip_summed = CHECKSUM_NONE;

	/* Remove ipcomp header and decompress original payload */
82
	iph = ipv6_hdr(skb);
83
	ipch = (void *)skb->data;
84
	skb->transport_header = skb->network_header + sizeof(*ipch);
85
	__skb_pull(skb, sizeof(*ipch));
L
Linus Torvalds 已提交
86 87 88 89 90 91 92 93 94 95 96 97 98 99 100 101 102 103 104 105 106 107 108 109 110 111

	/* decompression */
	plen = skb->len;
	dlen = IPCOMP_SCRATCH_SIZE;
	start = skb->data;

	cpu = get_cpu();
	scratch = *per_cpu_ptr(ipcomp6_scratches, cpu);
	tfm = *per_cpu_ptr(ipcd->tfms, cpu);

	err = crypto_comp_decompress(tfm, start, plen, scratch, &dlen);
	if (err) {
		err = -EINVAL;
		goto out_put_cpu;
	}

	if (dlen < (plen + sizeof(struct ipv6_comp_hdr))) {
		err = -EINVAL;
		goto out_put_cpu;
	}

	err = pskb_expand_head(skb, 0, dlen - plen, GFP_ATOMIC);
	if (err) {
		goto out_put_cpu;
	}

112 113
	skb->truesize += dlen - plen;
	__skb_put(skb, dlen - plen);
L
Linus Torvalds 已提交
114
	memcpy(skb->data, scratch, dlen);
115
	err = ipch->nexthdr;
L
Linus Torvalds 已提交
116 117 118 119 120 121 122 123 124 125 126 127 128 129 130

out_put_cpu:
	put_cpu();
out:
	return err;
}

static int ipcomp6_output(struct xfrm_state *x, struct sk_buff *skb)
{
	int err;
	struct ipv6hdr *top_iph;
	struct ipv6_comp_hdr *ipch;
	struct ipcomp_data *ipcd = x->data;
	int plen, dlen;
	u8 *start, *scratch;
131
	struct crypto_comp *tfm;
L
Linus Torvalds 已提交
132
	int cpu;
133
	int hdr_len = skb_transport_offset(skb);
L
Linus Torvalds 已提交
134 135 136 137 138 139

	/* check whether datagram len is larger than threshold */
	if ((skb->len - hdr_len) < ipcd->threshold) {
		goto out_ok;
	}

H
Herbert Xu 已提交
140
	if (skb_linearize_cow(skb))
L
Linus Torvalds 已提交
141 142 143 144 145
		goto out_ok;

	/* compression */
	plen = skb->len - hdr_len;
	dlen = IPCOMP_SCRATCH_SIZE;
146
	start = skb_transport_header(skb);
L
Linus Torvalds 已提交
147 148 149 150 151 152 153 154 155 156 157 158 159 160 161 162 163 164 165 166

	cpu = get_cpu();
	scratch = *per_cpu_ptr(ipcomp6_scratches, cpu);
	tfm = *per_cpu_ptr(ipcd->tfms, cpu);

	err = crypto_comp_compress(tfm, start, plen, scratch, &dlen);
	if (err || (dlen + sizeof(struct ipv6_comp_hdr)) >= plen) {
		put_cpu();
		goto out_ok;
	}
	memcpy(start + sizeof(struct ip_comp_hdr), scratch, dlen);
	put_cpu();
	pskb_trim(skb, hdr_len + dlen + sizeof(struct ip_comp_hdr));

	/* insert ipcomp header and replace datagram */
	top_iph = (struct ipv6hdr *)skb->data;

	top_iph->payload_len = htons(skb->len - sizeof(struct ipv6hdr));

	ipch = (struct ipv6_comp_hdr *)start;
167
	ipch->nexthdr = *skb_network_header(skb);
L
Linus Torvalds 已提交
168 169
	ipch->flags = 0;
	ipch->cpi = htons((u16 )ntohl(x->id.spi));
170
	*skb_network_header(skb) = IPPROTO_COMP;
L
Linus Torvalds 已提交
171 172 173 174 175 176

out_ok:
	return 0;
}

static void ipcomp6_err(struct sk_buff *skb, struct inet6_skb_parm *opt,
177
				int type, int code, int offset, __be32 info)
L
Linus Torvalds 已提交
178
{
179
	__be32 spi;
L
Linus Torvalds 已提交
180 181 182 183 184 185 186
	struct ipv6hdr *iph = (struct ipv6hdr*)skb->data;
	struct ipv6_comp_hdr *ipcomph = (struct ipv6_comp_hdr*)(skb->data+offset);
	struct xfrm_state *x;

	if (type != ICMPV6_DEST_UNREACH && type != ICMPV6_PKT_TOOBIG)
		return;

A
Alexey Dobriyan 已提交
187
	spi = htonl(ntohs(ipcomph->cpi));
L
Linus Torvalds 已提交
188 189 190 191
	x = xfrm_state_lookup((xfrm_address_t *)&iph->daddr, spi, IPPROTO_COMP, AF_INET6);
	if (!x)
		return;

J
Joe Perches 已提交
192
	printk(KERN_DEBUG "pmtu discovery on SA IPCOMP/%08x/" NIP6_FMT "\n",
L
Linus Torvalds 已提交
193 194 195 196 197 198 199
			spi, NIP6(iph->daddr));
	xfrm_state_put(x);
}

static struct xfrm_state *ipcomp6_tunnel_create(struct xfrm_state *x)
{
	struct xfrm_state *t = NULL;
D
Diego Beltrami 已提交
200
	u8 mode = XFRM_MODE_TUNNEL;
L
Linus Torvalds 已提交
201 202 203 204 205 206 207

	t = xfrm_state_alloc();
	if (!t)
		goto out;

	t->id.proto = IPPROTO_IPV6;
	t->id.spi = xfrm6_tunnel_alloc_spi((xfrm_address_t *)&x->props.saddr);
208 209 210
	if (!t->id.spi)
		goto error;

L
Linus Torvalds 已提交
211 212 213
	memcpy(t->id.daddr.a6, x->id.daddr.a6, sizeof(struct in6_addr));
	memcpy(&t->sel, &x->sel, sizeof(t->sel));
	t->props.family = AF_INET6;
D
Diego Beltrami 已提交
214 215 216
	if (x->props.mode == XFRM_MODE_BEET)
		mode = x->props.mode;
	t->props.mode = mode;
L
Linus Torvalds 已提交
217 218
	memcpy(t->props.saddr.a6, x->props.saddr.a6, sizeof(struct in6_addr));

H
Herbert Xu 已提交
219
	if (xfrm_init_state(t))
L
Linus Torvalds 已提交
220 221 222 223 224 225 226 227
		goto error;

	atomic_set(&t->tunnel_users, 1);

out:
	return t;

error:
228
	t->km.state = XFRM_STATE_DEAD;
L
Linus Torvalds 已提交
229
	xfrm_state_put(t);
230
	t = NULL;
L
Linus Torvalds 已提交
231 232 233 234 235 236 237
	goto out;
}

static int ipcomp6_tunnel_attach(struct xfrm_state *x)
{
	int err = 0;
	struct xfrm_state *t = NULL;
238
	__be32 spi;
L
Linus Torvalds 已提交
239 240 241 242 243 244 245 246 247 248 249 250 251 252 253 254 255 256 257 258 259 260 261 262 263 264 265 266 267 268 269 270 271

	spi = xfrm6_tunnel_spi_lookup((xfrm_address_t *)&x->props.saddr);
	if (spi)
		t = xfrm_state_lookup((xfrm_address_t *)&x->id.daddr,
					      spi, IPPROTO_IPV6, AF_INET6);
	if (!t) {
		t = ipcomp6_tunnel_create(x);
		if (!t) {
			err = -EINVAL;
			goto out;
		}
		xfrm_state_insert(t);
		xfrm_state_hold(t);
	}
	x->tunnel = t;
	atomic_inc(&t->tunnel_users);

out:
	return err;
}

static void ipcomp6_free_scratches(void)
{
	int i;
	void **scratches;

	if (--ipcomp6_scratch_users)
		return;

	scratches = ipcomp6_scratches;
	if (!scratches)
		return;

272
	for_each_possible_cpu(i) {
L
Linus Torvalds 已提交
273
		void *scratch = *per_cpu_ptr(scratches, i);
274 275

		vfree(scratch);
L
Linus Torvalds 已提交
276 277 278 279 280 281 282 283 284 285 286 287 288 289 290 291 292 293 294
	}

	free_percpu(scratches);
}

static void **ipcomp6_alloc_scratches(void)
{
	int i;
	void **scratches;

	if (ipcomp6_scratch_users++)
		return ipcomp6_scratches;

	scratches = alloc_percpu(void *);
	if (!scratches)
		return NULL;

	ipcomp6_scratches = scratches;

295
	for_each_possible_cpu(i) {
L
Linus Torvalds 已提交
296 297 298 299 300 301 302 303 304
		void *scratch = vmalloc(IPCOMP_SCRATCH_SIZE);
		if (!scratch)
			return NULL;
		*per_cpu_ptr(scratches, i) = scratch;
	}

	return scratches;
}

305
static void ipcomp6_free_tfms(struct crypto_comp **tfms)
L
Linus Torvalds 已提交
306 307 308 309 310 311 312 313 314 315 316 317 318 319 320 321 322 323 324 325
{
	struct ipcomp6_tfms *pos;
	int cpu;

	list_for_each_entry(pos, &ipcomp6_tfms_list, list) {
		if (pos->tfms == tfms)
			break;
	}

	BUG_TRAP(pos);

	if (--pos->users)
		return;

	list_del(&pos->list);
	kfree(pos);

	if (!tfms)
		return;

326
	for_each_possible_cpu(cpu) {
327 328
		struct crypto_comp *tfm = *per_cpu_ptr(tfms, cpu);
		crypto_free_comp(tfm);
L
Linus Torvalds 已提交
329 330 331 332
	}
	free_percpu(tfms);
}

333
static struct crypto_comp **ipcomp6_alloc_tfms(const char *alg_name)
L
Linus Torvalds 已提交
334 335
{
	struct ipcomp6_tfms *pos;
336
	struct crypto_comp **tfms;
L
Linus Torvalds 已提交
337 338 339
	int cpu;

	/* This can be any valid CPU ID so we don't need locking. */
340
	cpu = raw_smp_processor_id();
L
Linus Torvalds 已提交
341 342

	list_for_each_entry(pos, &ipcomp6_tfms_list, list) {
343
		struct crypto_comp *tfm;
L
Linus Torvalds 已提交
344 345 346 347

		tfms = pos->tfms;
		tfm = *per_cpu_ptr(tfms, cpu);

348
		if (!strcmp(crypto_comp_name(tfm), alg_name)) {
L
Linus Torvalds 已提交
349 350 351 352 353 354 355 356 357 358 359 360 361
			pos->users++;
			return tfms;
		}
	}

	pos = kmalloc(sizeof(*pos), GFP_KERNEL);
	if (!pos)
		return NULL;

	pos->users = 1;
	INIT_LIST_HEAD(&pos->list);
	list_add(&pos->list, &ipcomp6_tfms_list);

362
	pos->tfms = tfms = alloc_percpu(struct crypto_comp *);
L
Linus Torvalds 已提交
363 364 365
	if (!tfms)
		goto error;

366
	for_each_possible_cpu(cpu) {
367 368
		struct crypto_comp *tfm = crypto_alloc_comp(alg_name, 0,
							    CRYPTO_ALG_ASYNC);
L
Linus Torvalds 已提交
369 370 371 372 373 374 375 376 377 378 379 380 381 382 383 384 385 386 387 388 389 390 391 392 393
		if (!tfm)
			goto error;
		*per_cpu_ptr(tfms, cpu) = tfm;
	}

	return tfms;

error:
	ipcomp6_free_tfms(tfms);
	return NULL;
}

static void ipcomp6_free_data(struct ipcomp_data *ipcd)
{
	if (ipcd->tfms)
		ipcomp6_free_tfms(ipcd->tfms);
	ipcomp6_free_scratches();
}

static void ipcomp6_destroy(struct xfrm_state *x)
{
	struct ipcomp_data *ipcd = x->data;
	if (!ipcd)
		return;
	xfrm_state_delete_tunnel(x);
A
Arjan van de Ven 已提交
394
	mutex_lock(&ipcomp6_resource_mutex);
L
Linus Torvalds 已提交
395
	ipcomp6_free_data(ipcd);
A
Arjan van de Ven 已提交
396
	mutex_unlock(&ipcomp6_resource_mutex);
L
Linus Torvalds 已提交
397 398 399 400 401
	kfree(ipcd);

	xfrm6_tunnel_free_spi((xfrm_address_t *)&x->props.saddr);
}

H
Herbert Xu 已提交
402
static int ipcomp6_init_state(struct xfrm_state *x)
L
Linus Torvalds 已提交
403 404 405 406 407 408 409 410 411 412 413 414 415
{
	int err;
	struct ipcomp_data *ipcd;
	struct xfrm_algo_desc *calg_desc;

	err = -EINVAL;
	if (!x->calg)
		goto out;

	if (x->encap)
		goto out;

	err = -ENOMEM;
416
	ipcd = kzalloc(sizeof(*ipcd), GFP_KERNEL);
L
Linus Torvalds 已提交
417 418 419 420
	if (!ipcd)
		goto out;

	x->props.header_len = 0;
421
	if (x->props.mode == XFRM_MODE_TUNNEL)
L
Linus Torvalds 已提交
422
		x->props.header_len += sizeof(struct ipv6hdr);
423

A
Arjan van de Ven 已提交
424
	mutex_lock(&ipcomp6_resource_mutex);
L
Linus Torvalds 已提交
425 426 427 428 429 430
	if (!ipcomp6_alloc_scratches())
		goto error;

	ipcd->tfms = ipcomp6_alloc_tfms(x->calg->alg_name);
	if (!ipcd->tfms)
		goto error;
A
Arjan van de Ven 已提交
431
	mutex_unlock(&ipcomp6_resource_mutex);
L
Linus Torvalds 已提交
432

433
	if (x->props.mode == XFRM_MODE_TUNNEL) {
L
Linus Torvalds 已提交
434 435 436 437 438 439 440 441 442 443 444 445 446
		err = ipcomp6_tunnel_attach(x);
		if (err)
			goto error_tunnel;
	}

	calg_desc = xfrm_calg_get_byname(x->calg->alg_name, 0);
	BUG_ON(!calg_desc);
	ipcd->threshold = calg_desc->uinfo.comp.threshold;
	x->data = ipcd;
	err = 0;
out:
	return err;
error_tunnel:
A
Arjan van de Ven 已提交
447
	mutex_lock(&ipcomp6_resource_mutex);
L
Linus Torvalds 已提交
448 449
error:
	ipcomp6_free_data(ipcd);
A
Arjan van de Ven 已提交
450
	mutex_unlock(&ipcomp6_resource_mutex);
L
Linus Torvalds 已提交
451 452 453 454 455
	kfree(ipcd);

	goto out;
}

456
static struct xfrm_type ipcomp6_type =
L
Linus Torvalds 已提交
457 458 459 460 461 462 463 464
{
	.description	= "IPCOMP6",
	.owner		= THIS_MODULE,
	.proto		= IPPROTO_COMP,
	.init_state	= ipcomp6_init_state,
	.destructor	= ipcomp6_destroy,
	.input		= ipcomp6_input,
	.output		= ipcomp6_output,
465
	.hdr_offset	= xfrm6_find_1stfragopt,
L
Linus Torvalds 已提交
466 467
};

468
static struct inet6_protocol ipcomp6_protocol =
L
Linus Torvalds 已提交
469 470 471 472 473 474 475 476 477 478 479 480 481 482 483 484 485 486 487 488 489 490
{
	.handler	= xfrm6_rcv,
	.err_handler	= ipcomp6_err,
	.flags		= INET6_PROTO_NOPOLICY,
};

static int __init ipcomp6_init(void)
{
	if (xfrm_register_type(&ipcomp6_type, AF_INET6) < 0) {
		printk(KERN_INFO "ipcomp6 init: can't add xfrm type\n");
		return -EAGAIN;
	}
	if (inet6_add_protocol(&ipcomp6_protocol, IPPROTO_COMP) < 0) {
		printk(KERN_INFO "ipcomp6 init: can't add protocol\n");
		xfrm_unregister_type(&ipcomp6_type, AF_INET6);
		return -EAGAIN;
	}
	return 0;
}

static void __exit ipcomp6_fini(void)
{
491
	if (inet6_del_protocol(&ipcomp6_protocol, IPPROTO_COMP) < 0)
L
Linus Torvalds 已提交
492 493 494 495 496 497 498 499 500 501 502 503
		printk(KERN_INFO "ipv6 ipcomp close: can't remove protocol\n");
	if (xfrm_unregister_type(&ipcomp6_type, AF_INET6) < 0)
		printk(KERN_INFO "ipv6 ipcomp close: can't remove xfrm type\n");
}

module_init(ipcomp6_init);
module_exit(ipcomp6_fini);
MODULE_LICENSE("GPL");
MODULE_DESCRIPTION("IP Payload Compression Protocol (IPComp) for IPv6 - RFC3173");
MODULE_AUTHOR("Mitsuru KANDA <mk@linux-ipv6.org>");