tls.h 15.4 KB
Newer Older
D
Dave Watson 已提交
1 2 3 4 5 6 7 8 9 10 11 12 13 14 15 16 17 18 19 20 21 22 23 24 25 26 27 28 29 30 31 32 33 34 35 36 37
/*
 * Copyright (c) 2016-2017, Mellanox Technologies. All rights reserved.
 * Copyright (c) 2016-2017, Dave Watson <davejwatson@fb.com>. All rights reserved.
 *
 * This software is available to you under a choice of one of two
 * licenses.  You may choose to be licensed under the terms of the GNU
 * General Public License (GPL) Version 2, available from the file
 * COPYING in the main directory of this source tree, or the
 * OpenIB.org BSD license below:
 *
 *     Redistribution and use in source and binary forms, with or
 *     without modification, are permitted provided that the following
 *     conditions are met:
 *
 *      - Redistributions of source code must retain the above
 *        copyright notice, this list of conditions and the following
 *        disclaimer.
 *
 *      - Redistributions in binary form must reproduce the above
 *        copyright notice, this list of conditions and the following
 *        disclaimer in the documentation and/or other materials
 *        provided with the distribution.
 *
 * THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND,
 * EXPRESS OR IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF
 * MERCHANTABILITY, FITNESS FOR A PARTICULAR PURPOSE AND
 * NONINFRINGEMENT. IN NO EVENT SHALL THE AUTHORS OR COPYRIGHT HOLDERS
 * BE LIABLE FOR ANY CLAIM, DAMAGES OR OTHER LIABILITY, WHETHER IN AN
 * ACTION OF CONTRACT, TORT OR OTHERWISE, ARISING FROM, OUT OF OR IN
 * CONNECTION WITH THE SOFTWARE OR THE USE OR OTHER DEALINGS IN THE
 * SOFTWARE.
 */

#ifndef _TLS_OFFLOAD_H
#define _TLS_OFFLOAD_H

#include <linux/types.h>
38
#include <asm/byteorder.h>
39
#include <linux/crypto.h>
40 41
#include <linux/socket.h>
#include <linux/tcp.h>
42 43
#include <linux/skmsg.h>

44
#include <net/tcp.h>
D
Dave Watson 已提交
45
#include <net/strparser.h>
46
#include <crypto/aead.h>
D
Dave Watson 已提交
47 48 49 50 51 52 53 54 55 56 57 58 59 60
#include <uapi/linux/tls.h>


/* Maximum data size carried in a TLS record */
#define TLS_MAX_PAYLOAD_SIZE		((size_t)1 << 14)

#define TLS_HEADER_SIZE			5
#define TLS_NONCE_OFFSET		TLS_HEADER_SIZE

#define TLS_CRYPTO_INFO_READY(info)	((info)->cipher_type)

#define TLS_RECORD_TYPE_DATA		0x17

#define TLS_AAD_SPACE_SIZE		13
A
Atul Gupta 已提交
61 62 63 64 65 66 67 68 69 70 71 72 73 74 75 76 77 78
#define TLS_DEVICE_NAME_MAX		32

/*
 * This structure defines the routines for Inline TLS driver.
 * The following routines are optional and filled with a
 * null pointer if not defined.
 *
 * @name: Its the name of registered Inline tls device
 * @dev_list: Inline tls device list
 * int (*feature)(struct tls_device *device);
 *     Called to return Inline TLS driver capability
 *
 * int (*hash)(struct tls_device *device, struct sock *sk);
 *     This function sets Inline driver for listen and program
 *     device specific functioanlity as required
 *
 * void (*unhash)(struct tls_device *device, struct sock *sk);
 *     This function cleans listen state set by Inline TLS driver
79 80 81 82
 *
 * void (*release)(struct kref *kref);
 *     Release the registered device and allocated resources
 * @kref: Number of reference to tls_device
A
Atul Gupta 已提交
83 84 85 86 87 88 89
 */
struct tls_device {
	char name[TLS_DEVICE_NAME_MAX];
	struct list_head dev_list;
	int  (*feature)(struct tls_device *device);
	int  (*hash)(struct tls_device *device, struct sock *sk);
	void (*unhash)(struct tls_device *device, struct sock *sk);
90 91
	void (*release)(struct kref *kref);
	struct kref kref;
A
Atul Gupta 已提交
92
};
D
Dave Watson 已提交
93

94 95 96 97 98 99 100 101 102 103
enum {
	TLS_BASE,
	TLS_SW,
#ifdef CONFIG_TLS_DEVICE
	TLS_HW,
#endif
	TLS_HW_RECORD,
	TLS_NUM_CONFIG,
};

104 105 106 107 108 109
/* TLS records are maintained in 'struct tls_rec'. It stores the memory pages
 * allocated or mapped for each TLS record. After encryption, the records are
 * stores in a linked list.
 */
struct tls_rec {
	struct list_head list;
110
	int tx_ready;
111
	int tx_flags;
112
	int inplace_crypto;
D
Dave Watson 已提交
113

114 115
	struct sk_msg msg_plaintext;
	struct sk_msg msg_encrypted;
116

117 118 119 120
	/* AAD | msg_plaintext.sg.data | sg_tag */
	struct scatterlist sg_aead_in[2];
	/* AAD | msg_encrypted.sg.data (data contains overhead for hdr & iv & tag) */
	struct scatterlist sg_aead_out[2];
121

D
Dave Watson 已提交
122 123 124
	char content_type;
	struct scatterlist sg_content_type;

125
	char aad_space[TLS_AAD_SPACE_SIZE];
126 127
	u8 iv_data[TLS_CIPHER_AES_GCM_128_IV_SIZE +
		   TLS_CIPHER_AES_GCM_128_SALT_SIZE];
128 129 130 131
	struct aead_request aead_req;
	u8 aead_req_ctx[];
};

132 133 134 135 136
struct tls_msg {
	struct strp_msg rxm;
	u8 control;
};

137 138 139 140 141 142 143 144 145 146
struct tx_work {
	struct delayed_work work;
	struct sock *sk;
};

struct tls_sw_context_tx {
	struct crypto_aead *aead_send;
	struct crypto_wait async_wait;
	struct tx_work tx_work;
	struct tls_rec *open_rec;
147
	struct list_head tx_list;
148 149
	atomic_t encrypt_pending;
	int async_notify;
150
	int async_capable;
151 152 153

#define BIT_TX_SCHEDULED	0
	unsigned long tx_bitmask;
D
Dave Watson 已提交
154 155
};

B
Boris Pismenny 已提交
156 157 158 159
struct tls_sw_context_rx {
	struct crypto_aead *aead_recv;
	struct crypto_wait async_wait;
	struct strparser strp;
160
	struct sk_buff_head rx_list;	/* list of decrypted 'data' records */
B
Boris Pismenny 已提交
161
	void (*saved_data_ready)(struct sock *sk);
162

B
Boris Pismenny 已提交
163 164
	struct sk_buff *recv_pkt;
	u8 control;
165
	int async_capable;
B
Boris Pismenny 已提交
166
	bool decrypted;
167 168 169 170
	atomic_t decrypt_pending;
	bool async_notify;
};

171 172 173 174 175 176 177 178
struct tls_record_info {
	struct list_head list;
	u32 end_seq;
	int len;
	int num_frags;
	skb_frag_t frags[MAX_SKB_FRAGS];
};

179
struct tls_offload_context_tx {
180 181 182 183 184 185 186 187 188 189 190 191 192 193 194 195 196 197
	struct crypto_aead *aead_send;
	spinlock_t lock;	/* protects records list */
	struct list_head records_list;
	struct tls_record_info *open_record;
	struct tls_record_info *retransmit_hint;
	u64 hint_record_sn;
	u64 unacked_record_sn;

	struct scatterlist sg_tx_data[MAX_SKB_FRAGS];
	void (*sk_destruct)(struct sock *sk);
	u8 driver_state[];
	/* The TLS layer reserves room for driver specific state
	 * Currently the belief is that there is not enough
	 * driver specific state to justify another layer of indirection
	 */
#define TLS_DRIVER_STATE_SIZE (max_t(size_t, 8, sizeof(void *)))
};

198 199
#define TLS_OFFLOAD_CONTEXT_SIZE_TX                                            \
	(ALIGN(sizeof(struct tls_offload_context_tx), sizeof(void *)) +        \
200 201
	 TLS_DRIVER_STATE_SIZE)

D
Dave Watson 已提交
202 203 204 205
enum {
	TLS_PENDING_CLOSED_RECORD
};

206 207 208 209 210
struct cipher_context {
	char *iv;
	char *rec_seq;
};

211 212
union tls_crypto_context {
	struct tls_crypto_info info;
D
Dave Watson 已提交
213 214 215 216
	union {
		struct tls12_crypto_info_aes_gcm_128 aes_gcm_128;
		struct tls12_crypto_info_aes_gcm_256 aes_gcm_256;
	};
217 218
};

219 220 221 222 223 224 225 226 227 228 229 230
struct tls_prot_info {
	u16 version;
	u16 cipher_type;
	u16 prepend_size;
	u16 tag_size;
	u16 overhead_size;
	u16 iv_size;
	u16 rec_seq_size;
	u16 aad_size;
	u16 tail_size;
};

D
Dave Watson 已提交
231
struct tls_context {
232 233
	struct tls_prot_info prot_info;

234 235
	union tls_crypto_context crypto_send;
	union tls_crypto_context crypto_recv;
D
Dave Watson 已提交
236

B
Boris Pismenny 已提交
237 238 239 240 241 242
	struct list_head list;
	struct net_device *netdev;
	refcount_t refcount;

	void *priv_ctx_tx;
	void *priv_ctx_rx;
D
Dave Watson 已提交
243

B
Boris Pismenny 已提交
244 245
	u8 tx_conf:3;
	u8 rx_conf:3;
246

247
	struct cipher_context tx;
D
Dave Watson 已提交
248
	struct cipher_context rx;
D
Dave Watson 已提交
249 250 251

	struct scatterlist *partially_sent_record;
	u16 partially_sent_offset;
252

D
Dave Watson 已提交
253
	unsigned long flags;
254
	bool in_tcp_sendpages;
255
	bool pending_open_record_frags;
D
Dave Watson 已提交
256 257 258 259

	int (*push_pending_record)(struct sock *sk, int flags);

	void (*sk_write_space)(struct sock *sk);
260
	void (*sk_destruct)(struct sock *sk);
D
Dave Watson 已提交
261 262 263 264 265 266 267 268
	void (*sk_proto_close)(struct sock *sk, long timeout);

	int  (*setsockopt)(struct sock *sk, int level,
			   int optname, char __user *optval,
			   unsigned int optlen);
	int  (*getsockopt)(struct sock *sk, int level,
			   int optname, char __user *optval,
			   int __user *optlen);
A
Atul Gupta 已提交
269 270
	int  (*hash)(struct sock *sk);
	void (*unhash)(struct sock *sk);
D
Dave Watson 已提交
271 272
};

273 274 275 276 277 278 279 280 281 282 283 284 285 286 287
struct tls_offload_context_rx {
	/* sw must be the first member of tls_offload_context_rx */
	struct tls_sw_context_rx sw;
	atomic64_t resync_req;
	u8 driver_state[];
	/* The TLS layer reserves room for driver specific state
	 * Currently the belief is that there is not enough
	 * driver specific state to justify another layer of indirection
	 */
};

#define TLS_OFFLOAD_CONTEXT_SIZE_RX					\
	(ALIGN(sizeof(struct tls_offload_context_rx), sizeof(void *)) + \
	 TLS_DRIVER_STATE_SIZE)

D
Dave Watson 已提交
288 289 290 291 292 293
int wait_on_pending_writer(struct sock *sk, long *timeo);
int tls_sk_query(struct sock *sk, int optname, char __user *optval,
		int __user *optlen);
int tls_sk_attach(struct sock *sk, int optname, char __user *optval,
		  unsigned int optlen);

D
Dave Watson 已提交
294
int tls_set_sw_offload(struct sock *sk, struct tls_context *ctx, int tx);
D
Dave Watson 已提交
295 296 297 298
int tls_sw_sendmsg(struct sock *sk, struct msghdr *msg, size_t size);
int tls_sw_sendpage(struct sock *sk, struct page *page,
		    int offset, size_t size, int flags);
void tls_sw_close(struct sock *sk, long timeout);
B
Boris Pismenny 已提交
299 300
void tls_sw_free_resources_tx(struct sock *sk);
void tls_sw_free_resources_rx(struct sock *sk);
301
void tls_sw_release_resources_rx(struct sock *sk);
D
Dave Watson 已提交
302 303
int tls_sw_recvmsg(struct sock *sk, struct msghdr *msg, size_t len,
		   int nonblock, int flags, int *addr_len);
304
bool tls_sw_stream_read(const struct sock *sk);
D
Dave Watson 已提交
305 306 307
ssize_t tls_sw_splice_read(struct socket *sock, loff_t *ppos,
			   struct pipe_inode_info *pipe,
			   size_t len, unsigned int flags);
D
Dave Watson 已提交
308

309 310 311 312 313 314 315
int tls_set_device_offload(struct sock *sk, struct tls_context *ctx);
int tls_device_sendmsg(struct sock *sk, struct msghdr *msg, size_t size);
int tls_device_sendpage(struct sock *sk, struct page *page,
			int offset, size_t size, int flags);
void tls_device_sk_destruct(struct sock *sk);
void tls_device_init(void);
void tls_device_cleanup(void);
316
int tls_tx_records(struct sock *sk, int flags);
317

318
struct tls_record_info *tls_get_record(struct tls_offload_context_tx *context,
319 320 321 322 323 324 325 326 327 328 329
				       u32 seq, u64 *p_record_sn);

static inline bool tls_record_is_start_marker(struct tls_record_info *rec)
{
	return rec->len == 0;
}

static inline u32 tls_record_start_seq(struct tls_record_info *rec)
{
	return rec->end_seq - rec->len;
}
D
Dave Watson 已提交
330

331
void tls_sk_destruct(struct sock *sk, struct tls_context *ctx);
D
Dave Watson 已提交
332 333 334
int tls_push_sg(struct sock *sk, struct tls_context *ctx,
		struct scatterlist *sg, u16 first_offset,
		int flags);
335 336 337
int tls_push_partial_record(struct sock *sk, struct tls_context *ctx,
			    int flags);

D
Dave Watson 已提交
338 339 340
int tls_push_pending_closed_record(struct sock *sk, struct tls_context *ctx,
				   int flags, long *timeo);

341 342 343 344 345
static inline struct tls_msg *tls_msg(struct sk_buff *skb)
{
	return (struct tls_msg *)strp_msg(skb);
}

D
Dave Watson 已提交
346 347 348 349 350 351 352 353 354 355 356 357 358 359 360 361 362 363 364 365 366 367 368 369 370 371 372 373 374 375
static inline bool tls_is_pending_closed_record(struct tls_context *ctx)
{
	return test_bit(TLS_PENDING_CLOSED_RECORD, &ctx->flags);
}

static inline int tls_complete_pending_work(struct sock *sk,
					    struct tls_context *ctx,
					    int flags, long *timeo)
{
	int rc = 0;

	if (unlikely(sk->sk_write_pending))
		rc = wait_on_pending_writer(sk, timeo);

	if (!rc && tls_is_pending_closed_record(ctx))
		rc = tls_push_pending_closed_record(sk, ctx, flags, timeo);

	return rc;
}

static inline bool tls_is_partially_sent_record(struct tls_context *ctx)
{
	return !!ctx->partially_sent_record;
}

static inline bool tls_is_pending_open_record(struct tls_context *tls_ctx)
{
	return tls_ctx->pending_open_record_frags;
}

376
static inline bool is_tx_ready(struct tls_sw_context_tx *ctx)
377 378 379
{
	struct tls_rec *rec;

380
	rec = list_first_entry(&ctx->tx_list, struct tls_rec, list);
381 382 383
	if (!rec)
		return false;

384
	return READ_ONCE(rec->tx_ready);
385 386
}

387 388 389 390
struct sk_buff *
tls_validate_xmit_skb(struct sock *sk, struct net_device *dev,
		      struct sk_buff *skb);

391 392
static inline bool tls_is_sk_tx_device_offloaded(struct sock *sk)
{
393 394 395 396 397 398 399
#ifdef CONFIG_SOCK_VALIDATE_XMIT
	return sk_fullsock(sk) &
	       (smp_load_acquire(&sk->sk_validate_xmit_skb) ==
	       &tls_validate_xmit_skb);
#else
	return false;
#endif
400 401
}

402
static inline void tls_err_abort(struct sock *sk, int err)
D
Dave Watson 已提交
403
{
404
	sk->sk_err = err;
D
Dave Watson 已提交
405 406 407 408 409 410 411 412 413 414 415 416 417 418 419 420
	sk->sk_error_report(sk);
}

static inline bool tls_bigint_increment(unsigned char *seq, int len)
{
	int i;

	for (i = len - 1; i >= 0; i--) {
		++seq[i];
		if (seq[i] != 0)
			break;
	}

	return (i == -1);
}

421 422 423 424 425 426 427
static inline struct tls_context *tls_get_ctx(const struct sock *sk)
{
	struct inet_connection_sock *icsk = inet_csk(sk);

	return icsk->icsk_ulp_data;
}

D
Dave Watson 已提交
428
static inline void tls_advance_record_sn(struct sock *sk,
D
Dave Watson 已提交
429 430
					 struct cipher_context *ctx,
					 int version)
D
Dave Watson 已提交
431
{
432 433 434 435
	struct tls_context *tls_ctx = tls_get_ctx(sk);
	struct tls_prot_info *prot = &tls_ctx->prot_info;

	if (tls_bigint_increment(ctx->rec_seq, prot->rec_seq_size))
436
		tls_err_abort(sk, EBADMSG);
D
Dave Watson 已提交
437 438 439

	if (version != TLS_1_3_VERSION) {
		tls_bigint_increment(ctx->iv + TLS_CIPHER_AES_GCM_128_SALT_SIZE,
440
				     prot->iv_size);
D
Dave Watson 已提交
441
	}
D
Dave Watson 已提交
442 443 444 445 446
}

static inline void tls_fill_prepend(struct tls_context *ctx,
			     char *buf,
			     size_t plaintext_len,
D
Dave Watson 已提交
447 448
			     unsigned char record_type,
			     int version)
D
Dave Watson 已提交
449
{
450 451
	struct tls_prot_info *prot = &ctx->prot_info;
	size_t pkt_len, iv_size = prot->iv_size;
D
Dave Watson 已提交
452

453
	pkt_len = plaintext_len + prot->tag_size;
D
Dave Watson 已提交
454 455 456 457 458 459
	if (version != TLS_1_3_VERSION) {
		pkt_len += iv_size;

		memcpy(buf + TLS_NONCE_OFFSET,
		       ctx->tx.iv + TLS_CIPHER_AES_GCM_128_SALT_SIZE, iv_size);
	}
D
Dave Watson 已提交
460 461 462 463

	/* we cover nonce explicit here as well, so buf should be of
	 * size KTLS_DTLS_HEADER_SIZE + KTLS_DTLS_NONCE_EXPLICIT_SIZE
	 */
D
Dave Watson 已提交
464 465 466 467 468
	buf[0] = version == TLS_1_3_VERSION ?
		   TLS_RECORD_TYPE_DATA : record_type;
	/* Note that VERSION must be TLS_1_2 for both TLS1.2 and TLS1.3 */
	buf[1] = TLS_1_2_VERSION_MINOR;
	buf[2] = TLS_1_2_VERSION_MAJOR;
D
Dave Watson 已提交
469 470 471 472 473
	/* we can use IV for nonce explicit according to spec */
	buf[3] = pkt_len >> 8;
	buf[4] = pkt_len & 0xFF;
}

474 475 476 477
static inline void tls_make_aad(char *buf,
				size_t size,
				char *record_sequence,
				int record_sequence_size,
D
Dave Watson 已提交
478 479 480 481 482 483 484 485 486 487 488 489 490 491 492 493 494 495 496
				unsigned char record_type,
				int version)
{
	if (version != TLS_1_3_VERSION) {
		memcpy(buf, record_sequence, record_sequence_size);
		buf += 8;
	} else {
		size += TLS_CIPHER_AES_GCM_128_TAG_SIZE;
	}

	buf[0] = version == TLS_1_3_VERSION ?
		  TLS_RECORD_TYPE_DATA : record_type;
	buf[1] = TLS_1_2_VERSION_MAJOR;
	buf[2] = TLS_1_2_VERSION_MINOR;
	buf[3] = size >> 8;
	buf[4] = size & 0xFF;
}

static inline void xor_iv_with_seq(int version, char *iv, char *seq)
497
{
D
Dave Watson 已提交
498
	int i;
499

D
Dave Watson 已提交
500 501 502 503
	if (version == TLS_1_3_VERSION) {
		for (i = 0; i < 8; i++)
			iv[i + 4] ^= seq[i];
	}
504 505
}

D
Dave Watson 已提交
506

B
Boris Pismenny 已提交
507 508 509 510 511 512 513
static inline struct tls_sw_context_rx *tls_sw_ctx_rx(
		const struct tls_context *tls_ctx)
{
	return (struct tls_sw_context_rx *)tls_ctx->priv_ctx_rx;
}

static inline struct tls_sw_context_tx *tls_sw_ctx_tx(
D
Dave Watson 已提交
514 515
		const struct tls_context *tls_ctx)
{
B
Boris Pismenny 已提交
516
	return (struct tls_sw_context_tx *)tls_ctx->priv_ctx_tx;
D
Dave Watson 已提交
517 518
}

519 520
static inline struct tls_offload_context_tx *
tls_offload_ctx_tx(const struct tls_context *tls_ctx)
D
Dave Watson 已提交
521
{
522
	return (struct tls_offload_context_tx *)tls_ctx->priv_ctx_tx;
D
Dave Watson 已提交
523 524
}

525 526 527 528 529 530 531 532 533
static inline bool tls_sw_has_ctx_tx(const struct sock *sk)
{
	struct tls_context *ctx = tls_get_ctx(sk);

	if (!ctx)
		return false;
	return !!tls_sw_ctx_tx(ctx);
}

534 535 536 537 538 539 540 541 542 543 544 545 546 547 548 549
static inline struct tls_offload_context_rx *
tls_offload_ctx_rx(const struct tls_context *tls_ctx)
{
	return (struct tls_offload_context_rx *)tls_ctx->priv_ctx_rx;
}

/* The TLS context is valid until sk_destruct is called */
static inline void tls_offload_rx_resync_request(struct sock *sk, __be32 seq)
{
	struct tls_context *tls_ctx = tls_get_ctx(sk);
	struct tls_offload_context_rx *rx_ctx = tls_offload_ctx_rx(tls_ctx);

	atomic64_set(&rx_ctx->resync_req, ((((uint64_t)seq) << 32) | 1));
}


D
Dave Watson 已提交
550 551
int tls_proccess_cmsg(struct sock *sk, struct msghdr *msg,
		      unsigned char *record_type);
A
Atul Gupta 已提交
552 553
void tls_register_device(struct tls_device *device);
void tls_unregister_device(struct tls_device *device);
554
int tls_device_decrypted(struct sock *sk, struct sk_buff *skb);
555 556
int decrypt_skb(struct sock *sk, struct sk_buff *skb,
		struct scatterlist *sgout);
D
Dave Watson 已提交
557

558 559 560 561 562
struct sk_buff *tls_validate_xmit_skb(struct sock *sk,
				      struct net_device *dev,
				      struct sk_buff *skb);

int tls_sw_fallback_init(struct sock *sk,
563
			 struct tls_offload_context_tx *offload_ctx,
564 565
			 struct tls_crypto_info *crypto_info);

566 567 568 569 570
int tls_set_device_offload_rx(struct sock *sk, struct tls_context *ctx);

void tls_device_offload_cleanup_rx(struct sock *sk);
void handle_device_resync(struct sock *sk, u32 seq, u64 rcd_sn);

D
Dave Watson 已提交
571
#endif /* _TLS_OFFLOAD_H */