capability.c 9.9 KB
Newer Older
L
Linus Torvalds 已提交
1 2 3 4 5
/*
 * linux/kernel/capability.c
 *
 * Copyright (C) 1997  Andrew Main <zefram@fysh.org>
 *
6
 * Integrated into 2.1.97+,  Andrew G. Morgan <morgan@kernel.org>
L
Linus Torvalds 已提交
7
 * 30 May 2002:	Cleanup, Robert M. Love <rml@tech9.net>
8
 */
L
Linus Torvalds 已提交
9

10
#include <linux/capability.h>
L
Linus Torvalds 已提交
11 12 13 14
#include <linux/mm.h>
#include <linux/module.h>
#include <linux/security.h>
#include <linux/syscalls.h>
15
#include <linux/pid_namespace.h>
L
Linus Torvalds 已提交
16 17 18 19 20 21 22 23
#include <asm/uaccess.h>

/*
 * This lock protects task->cap_* for all tasks including current.
 * Locking rule: acquire this prior to tasklist_lock.
 */
static DEFINE_SPINLOCK(task_capability_lock);

24 25 26 27 28 29 30 31 32 33 34 35 36 37 38 39 40 41 42 43 44 45 46 47 48 49 50 51 52 53 54
/*
 * Leveraged for setting/resetting capabilities
 */

const kernel_cap_t __cap_empty_set = CAP_EMPTY_SET;
const kernel_cap_t __cap_full_set = CAP_FULL_SET;
const kernel_cap_t __cap_init_eff_set = CAP_INIT_EFF_SET;

EXPORT_SYMBOL(__cap_empty_set);
EXPORT_SYMBOL(__cap_full_set);
EXPORT_SYMBOL(__cap_init_eff_set);

/*
 * More recent versions of libcap are available from:
 *
 *   http://www.kernel.org/pub/linux/libs/security/linux-privs/
 */

static void warn_legacy_capability_use(void)
{
	static int warned;
	if (!warned) {
		char name[sizeof(current->comm)];

		printk(KERN_INFO "warning: `%s' uses 32-bit capabilities"
		       " (legacy support in use)\n",
		       get_task_comm(name, current));
		warned = 1;
	}
}

55 56 57 58 59 60 61 62 63 64 65 66 67 68 69 70 71 72 73 74 75 76 77 78 79 80 81 82 83 84 85 86 87 88 89 90 91 92 93 94 95 96 97 98 99 100 101 102 103 104 105 106 107 108 109 110 111 112 113 114 115 116 117
/*
 * Version 2 capabilities worked fine, but the linux/capability.h file
 * that accompanied their introduction encouraged their use without
 * the necessary user-space source code changes. As such, we have
 * created a version 3 with equivalent functionality to version 2, but
 * with a header change to protect legacy source code from using
 * version 2 when it wanted to use version 1. If your system has code
 * that trips the following warning, it is using version 2 specific
 * capabilities and may be doing so insecurely.
 *
 * The remedy is to either upgrade your version of libcap (to 2.10+,
 * if the application is linked against it), or recompile your
 * application with modern kernel headers and this warning will go
 * away.
 */

static void warn_deprecated_v2(void)
{
	static int warned;

	if (!warned) {
		char name[sizeof(current->comm)];

		printk(KERN_INFO "warning: `%s' uses deprecated v2"
		       " capabilities in a way that may be insecure.\n",
		       get_task_comm(name, current));
		warned = 1;
	}
}

/*
 * Version check. Return the number of u32s in each capability flag
 * array, or a negative value on error.
 */
static int cap_validate_magic(cap_user_header_t header, unsigned *tocopy)
{
	__u32 version;

	if (get_user(version, &header->version))
		return -EFAULT;

	switch (version) {
	case _LINUX_CAPABILITY_VERSION_1:
		warn_legacy_capability_use();
		*tocopy = _LINUX_CAPABILITY_U32S_1;
		break;
	case _LINUX_CAPABILITY_VERSION_2:
		warn_deprecated_v2();
		/*
		 * fall through - v3 is otherwise equivalent to v2.
		 */
	case _LINUX_CAPABILITY_VERSION_3:
		*tocopy = _LINUX_CAPABILITY_U32S_3;
		break;
	default:
		if (put_user((u32)_KERNEL_CAPABILITY_VERSION, &header->version))
			return -EFAULT;
		return -EINVAL;
	}

	return 0;
}

L
Linus Torvalds 已提交
118 119 120 121 122 123
/*
 * For sys_getproccap() and sys_setproccap(), any of the three
 * capability set pointers may be NULL -- indicating that that set is
 * uninteresting and/or not to be changed.
 */

124 125 126 127 128 129 130 131 132 133 134 135 136 137 138 139 140 141 142 143 144
/*
 * Atomically modify the effective capabilities returning the original
 * value. No permission check is performed here - it is assumed that the
 * caller is permitted to set the desired effective capabilities.
 */
kernel_cap_t cap_set_effective(const kernel_cap_t pE_new)
{
	kernel_cap_t pE_old;

	spin_lock(&task_capability_lock);

	pE_old = current->cap_effective;
	current->cap_effective = pE_new;

	spin_unlock(&task_capability_lock);

	return pE_old;
}

EXPORT_SYMBOL(cap_set_effective);

145
/**
L
Linus Torvalds 已提交
146
 * sys_capget - get the capabilities of a given process.
147 148 149 150 151 152
 * @header: pointer to struct that contains capability version and
 *	target pid data
 * @dataptr: pointer to struct that contains the effective, permitted,
 *	and inheritable capabilities that are returned
 *
 * Returns 0 on success and < 0 on error.
L
Linus Torvalds 已提交
153 154 155
 */
asmlinkage long sys_capget(cap_user_header_t header, cap_user_data_t dataptr)
{
156 157 158
	int ret = 0;
	pid_t pid;
	struct task_struct *target;
159 160
	unsigned tocopy;
	kernel_cap_t pE, pI, pP;
161

162 163 164
	ret = cap_validate_magic(header, &tocopy);
	if (ret != 0)
		return ret;
L
Linus Torvalds 已提交
165

166 167
	if (get_user(pid, &header->pid))
		return -EFAULT;
L
Linus Torvalds 已提交
168

169 170
	if (pid < 0)
		return -EINVAL;
L
Linus Torvalds 已提交
171

172 173
	spin_lock(&task_capability_lock);
	read_lock(&tasklist_lock);
L
Linus Torvalds 已提交
174

175
	if (pid && pid != task_pid_vnr(current)) {
176
		target = find_task_by_vpid(pid);
177 178 179 180 181 182
		if (!target) {
			ret = -ESRCH;
			goto out;
		}
	} else
		target = current;
L
Linus Torvalds 已提交
183

184
	ret = security_capget(target, &pE, &pI, &pP);
L
Linus Torvalds 已提交
185 186

out:
187 188
	read_unlock(&tasklist_lock);
	spin_unlock(&task_capability_lock);
L
Linus Torvalds 已提交
189

190
	if (!ret) {
191
		struct __user_cap_data_struct kdata[_KERNEL_CAPABILITY_U32S];
192 193 194 195 196 197 198 199 200
		unsigned i;

		for (i = 0; i < tocopy; i++) {
			kdata[i].effective = pE.cap[i];
			kdata[i].permitted = pP.cap[i];
			kdata[i].inheritable = pI.cap[i];
		}

		/*
201
		 * Note, in the case, tocopy < _KERNEL_CAPABILITY_U32S,
202 203 204 205 206 207 208 209 210 211 212 213 214 215 216 217 218 219 220 221 222 223 224
		 * we silently drop the upper capabilities here. This
		 * has the effect of making older libcap
		 * implementations implicitly drop upper capability
		 * bits when they perform a: capget/modify/capset
		 * sequence.
		 *
		 * This behavior is considered fail-safe
		 * behavior. Upgrading the application to a newer
		 * version of libcap will enable access to the newer
		 * capabilities.
		 *
		 * An alternative would be to return an error here
		 * (-ERANGE), but that causes legacy applications to
		 * unexpectidly fail; the capget/modify/capset aborts
		 * before modification is attempted and the application
		 * fails.
		 */

		if (copy_to_user(dataptr, kdata, tocopy
				 * sizeof(struct __user_cap_data_struct))) {
			return -EFAULT;
		}
	}
L
Linus Torvalds 已提交
225

226
	return ret;
L
Linus Torvalds 已提交
227 228 229 230 231 232
}

/*
 * cap_set_pg - set capabilities for all processes in a given process
 * group.  We call this holding task_capability_lock and tasklist_lock.
 */
233
static inline int cap_set_pg(int pgrp_nr, kernel_cap_t *effective,
L
Linus Torvalds 已提交
234 235 236
			      kernel_cap_t *inheritable,
			      kernel_cap_t *permitted)
{
237
	struct task_struct *g, *target;
L
Linus Torvalds 已提交
238 239
	int ret = -EPERM;
	int found = 0;
240
	struct pid *pgrp;
L
Linus Torvalds 已提交
241

242
	pgrp = find_vpid(pgrp_nr);
243
	do_each_pid_task(pgrp, PIDTYPE_PGID, g) {
L
Linus Torvalds 已提交
244 245 246 247 248 249 250 251 252 253 254 255
		target = g;
		while_each_thread(g, target) {
			if (!security_capset_check(target, effective,
							inheritable,
							permitted)) {
				security_capset_set(target, effective,
							inheritable,
							permitted);
				ret = 0;
			}
			found = 1;
		}
256
	} while_each_pid_task(pgrp, PIDTYPE_PGID, g);
L
Linus Torvalds 已提交
257 258

	if (!found)
259
		ret = 0;
L
Linus Torvalds 已提交
260 261 262 263 264 265 266 267 268 269 270
	return ret;
}

/*
 * cap_set_all - set capabilities for all processes other than init
 * and self.  We call this holding task_capability_lock and tasklist_lock.
 */
static inline int cap_set_all(kernel_cap_t *effective,
			       kernel_cap_t *inheritable,
			       kernel_cap_t *permitted)
{
271
     struct task_struct *g, *target;
L
Linus Torvalds 已提交
272 273 274 275
     int ret = -EPERM;
     int found = 0;

     do_each_thread(g, target) {
276
             if (target == current || is_container_init(target->group_leader))
L
Linus Torvalds 已提交
277 278 279 280 281 282 283 284 285 286 287 288 289 290
                     continue;
             found = 1;
	     if (security_capset_check(target, effective, inheritable,
						permitted))
		     continue;
	     ret = 0;
	     security_capset_set(target, effective, inheritable, permitted);
     } while_each_thread(g, target);

     if (!found)
	     ret = 0;
     return ret;
}

291 292 293 294 295 296 297 298
/**
 * sys_capset - set capabilities for a process or a group of processes
 * @header: pointer to struct that contains capability version and
 *	target pid data
 * @data: pointer to struct that contains the effective, permitted,
 *	and inheritable capabilities
 *
 * Set capabilities for a given process, all processes, or all
L
Linus Torvalds 已提交
299 300 301 302 303 304 305 306 307
 * processes in a given process group.
 *
 * The restrictions on setting capabilities are specified as:
 *
 * [pid is for the 'target' task.  'current' is the calling task.]
 *
 * I: any raised capabilities must be a subset of the (old current) permitted
 * P: any raised capabilities must be a subset of the (old current) permitted
 * E: must be set to a subset of (new target) permitted
308 309
 *
 * Returns 0 on success and < 0 on error.
L
Linus Torvalds 已提交
310 311 312
 */
asmlinkage long sys_capset(cap_user_header_t header, const cap_user_data_t data)
{
313
	struct __user_cap_data_struct kdata[_KERNEL_CAPABILITY_U32S];
314
	unsigned i, tocopy;
315 316 317 318 319
	kernel_cap_t inheritable, permitted, effective;
	struct task_struct *target;
	int ret;
	pid_t pid;

320 321 322
	ret = cap_validate_magic(header, &tocopy);
	if (ret != 0)
		return ret;
323 324 325 326

	if (get_user(pid, &header->pid))
		return -EFAULT;

327
	if (pid && pid != task_pid_vnr(current) && !capable(CAP_SETPCAP))
328 329
		return -EPERM;

330 331
	if (copy_from_user(&kdata, data, tocopy
			   * sizeof(struct __user_cap_data_struct))) {
332
		return -EFAULT;
333 334 335 336 337 338 339
	}

	for (i = 0; i < tocopy; i++) {
		effective.cap[i] = kdata[i].effective;
		permitted.cap[i] = kdata[i].permitted;
		inheritable.cap[i] = kdata[i].inheritable;
	}
340
	while (i < _KERNEL_CAPABILITY_U32S) {
341 342 343 344 345
		effective.cap[i] = 0;
		permitted.cap[i] = 0;
		inheritable.cap[i] = 0;
		i++;
	}
346 347 348 349

	spin_lock(&task_capability_lock);
	read_lock(&tasklist_lock);

350
	if (pid > 0 && pid != task_pid_vnr(current)) {
351
		target = find_task_by_vpid(pid);
352 353 354 355 356 357 358 359 360 361 362 363 364 365 366 367 368 369 370 371 372 373 374 375 376
		if (!target) {
			ret = -ESRCH;
			goto out;
		}
	} else
		target = current;

	ret = 0;

	/* having verified that the proposed changes are legal,
	   we now put them into effect. */
	if (pid < 0) {
		if (pid == -1)	/* all procs other than current and init */
			ret = cap_set_all(&effective, &inheritable, &permitted);

		else		/* all procs in process group */
			ret = cap_set_pg(-pid, &effective, &inheritable,
					 &permitted);
	} else {
		ret = security_capset_check(target, &effective, &inheritable,
					    &permitted);
		if (!ret)
			security_capset_set(target, &effective, &inheritable,
					    &permitted);
	}
L
Linus Torvalds 已提交
377 378

out:
379 380
	read_unlock(&tasklist_lock);
	spin_unlock(&task_capability_lock);
L
Linus Torvalds 已提交
381

382
	return ret;
L
Linus Torvalds 已提交
383
}
384 385 386 387 388 389 390 391 392 393 394 395 396 397 398

int __capable(struct task_struct *t, int cap)
{
	if (security_capable(t, cap) == 0) {
		t->flags |= PF_SUPERPRIV;
		return 1;
	}
	return 0;
}

int capable(int cap)
{
	return __capable(current, cap);
}
EXPORT_SYMBOL(capable);