dir.c 22.0 KB
Newer Older
L
Linus Torvalds 已提交
1 2 3 4 5 6 7 8 9 10
/*
 * dir.c - Operations for sysfs directories.
 */

#undef DEBUG

#include <linux/fs.h>
#include <linux/mount.h>
#include <linux/module.h>
#include <linux/kobject.h>
11
#include <linux/namei.h>
12
#include <linux/idr.h>
13
#include <linux/completion.h>
14
#include <asm/semaphore.h>
L
Linus Torvalds 已提交
15 16
#include "sysfs.h"

17
DEFINE_MUTEX(sysfs_mutex);
T
Tejun Heo 已提交
18
spinlock_t sysfs_assoc_lock = SPIN_LOCK_UNLOCKED;
L
Linus Torvalds 已提交
19

20 21 22
static spinlock_t sysfs_ino_lock = SPIN_LOCK_UNLOCKED;
static DEFINE_IDA(sysfs_ino_ida);

23 24 25 26 27 28 29 30
/**
 *	sysfs_link_sibling - link sysfs_dirent into sibling list
 *	@sd: sysfs_dirent of interest
 *
 *	Link @sd into its sibling list which starts from
 *	sd->s_parent->s_children.
 *
 *	Locking:
31
 *	mutex_lock(sysfs_mutex)
32 33 34 35 36 37 38 39 40 41 42 43 44 45 46 47 48 49
 */
static void sysfs_link_sibling(struct sysfs_dirent *sd)
{
	struct sysfs_dirent *parent_sd = sd->s_parent;

	BUG_ON(sd->s_sibling);
	sd->s_sibling = parent_sd->s_children;
	parent_sd->s_children = sd;
}

/**
 *	sysfs_unlink_sibling - unlink sysfs_dirent from sibling list
 *	@sd: sysfs_dirent of interest
 *
 *	Unlink @sd from its sibling list which starts from
 *	sd->s_parent->s_children.
 *
 *	Locking:
50
 *	mutex_lock(sysfs_mutex)
51 52 53 54 55 56 57 58 59 60 61 62 63 64
 */
static void sysfs_unlink_sibling(struct sysfs_dirent *sd)
{
	struct sysfs_dirent **pos;

	for (pos = &sd->s_parent->s_children; *pos; pos = &(*pos)->s_sibling) {
		if (*pos == sd) {
			*pos = sd->s_sibling;
			sd->s_sibling = NULL;
			break;
		}
	}
}

65 66 67 68 69 70 71 72 73 74 75 76
/**
 *	sysfs_get_active - get an active reference to sysfs_dirent
 *	@sd: sysfs_dirent to get an active reference to
 *
 *	Get an active reference of @sd.  This function is noop if @sd
 *	is NULL.
 *
 *	RETURNS:
 *	Pointer to @sd on success, NULL on failure.
 */
struct sysfs_dirent *sysfs_get_active(struct sysfs_dirent *sd)
{
77 78 79 80 81 82 83 84 85 86 87 88 89 90 91 92 93
	if (unlikely(!sd))
		return NULL;

	while (1) {
		int v, t;

		v = atomic_read(&sd->s_active);
		if (unlikely(v < 0))
			return NULL;

		t = atomic_cmpxchg(&sd->s_active, v, v + 1);
		if (likely(t == v))
			return sd;
		if (t < 0)
			return NULL;

		cpu_relax();
94 95 96 97 98 99 100 101 102 103 104 105
	}
}

/**
 *	sysfs_put_active - put an active reference to sysfs_dirent
 *	@sd: sysfs_dirent to put an active reference to
 *
 *	Put an active reference to @sd.  This function is noop if @sd
 *	is NULL.
 */
void sysfs_put_active(struct sysfs_dirent *sd)
{
106 107 108 109 110 111 112 113 114 115 116
	struct completion *cmpl;
	int v;

	if (unlikely(!sd))
		return;

	v = atomic_dec_return(&sd->s_active);
	if (likely(v != SD_DEACTIVATED_BIAS))
		return;

	/* atomic_dec_return() is a mb(), we'll always see the updated
117
	 * sd->s_sibling.
118
	 */
119
	cmpl = (void *)sd->s_sibling;
120
	complete(cmpl);
121 122 123 124 125 126 127 128 129 130 131 132 133 134 135 136 137 138 139 140 141 142 143 144 145 146 147 148 149 150 151 152 153 154 155 156 157 158 159 160 161 162 163 164 165
}

/**
 *	sysfs_get_active_two - get active references to sysfs_dirent and parent
 *	@sd: sysfs_dirent of interest
 *
 *	Get active reference to @sd and its parent.  Parent's active
 *	reference is grabbed first.  This function is noop if @sd is
 *	NULL.
 *
 *	RETURNS:
 *	Pointer to @sd on success, NULL on failure.
 */
struct sysfs_dirent *sysfs_get_active_two(struct sysfs_dirent *sd)
{
	if (sd) {
		if (sd->s_parent && unlikely(!sysfs_get_active(sd->s_parent)))
			return NULL;
		if (unlikely(!sysfs_get_active(sd))) {
			sysfs_put_active(sd->s_parent);
			return NULL;
		}
	}
	return sd;
}

/**
 *	sysfs_put_active_two - put active references to sysfs_dirent and parent
 *	@sd: sysfs_dirent of interest
 *
 *	Put active references to @sd and its parent.  This function is
 *	noop if @sd is NULL.
 */
void sysfs_put_active_two(struct sysfs_dirent *sd)
{
	if (sd) {
		sysfs_put_active(sd);
		sysfs_put_active(sd->s_parent);
	}
}

/**
 *	sysfs_deactivate - deactivate sysfs_dirent
 *	@sd: sysfs_dirent to deactivate
 *
166
 *	Deny new active references and drain existing ones.
167 168 169
 */
void sysfs_deactivate(struct sysfs_dirent *sd)
{
170 171
	DECLARE_COMPLETION_ONSTACK(wait);
	int v;
172

173
	BUG_ON(sd->s_sibling || !(sd->s_flags & SYSFS_FLAG_REMOVED));
174
	sd->s_sibling = (void *)&wait;
175 176

	/* atomic_add_return() is a mb(), put_active() will always see
177
	 * the updated sd->s_sibling.
178
	 */
179 180 181 182 183
	v = atomic_add_return(SD_DEACTIVATED_BIAS, &sd->s_active);

	if (v != SD_DEACTIVATED_BIAS)
		wait_for_completion(&wait);

184
	sd->s_sibling = NULL;
185 186
}

T
Tejun Heo 已提交
187
static int sysfs_alloc_ino(ino_t *pino)
188 189 190 191 192 193 194 195 196 197 198 199 200 201 202 203 204 205 206 207 208 209 210 211 212
{
	int ino, rc;

 retry:
	spin_lock(&sysfs_ino_lock);
	rc = ida_get_new_above(&sysfs_ino_ida, 2, &ino);
	spin_unlock(&sysfs_ino_lock);

	if (rc == -EAGAIN) {
		if (ida_pre_get(&sysfs_ino_ida, GFP_KERNEL))
			goto retry;
		rc = -ENOMEM;
	}

	*pino = ino;
	return rc;
}

static void sysfs_free_ino(ino_t ino)
{
	spin_lock(&sysfs_ino_lock);
	ida_remove(&sysfs_ino_ida, ino);
	spin_unlock(&sysfs_ino_lock);
}

213 214
void release_sysfs_dirent(struct sysfs_dirent * sd)
{
T
Tejun Heo 已提交
215 216 217
	struct sysfs_dirent *parent_sd;

 repeat:
218 219 220
	/* Moving/renaming is always done while holding reference.
	 * sd->s_parent won't change beneath us.
	 */
T
Tejun Heo 已提交
221 222
	parent_sd = sd->s_parent;

223
	if (sysfs_type(sd) == SYSFS_KOBJ_LINK)
224
		sysfs_put(sd->s_elem.symlink.target_sd);
225
	if (sysfs_type(sd) & SYSFS_COPY_NAME)
T
Tejun Heo 已提交
226
		kfree(sd->s_name);
227
	kfree(sd->s_iattr);
228
	sysfs_free_ino(sd->s_ino);
229
	kmem_cache_free(sysfs_dir_cachep, sd);
T
Tejun Heo 已提交
230 231 232 233

	sd = parent_sd;
	if (sd && atomic_dec_and_test(&sd->s_count))
		goto repeat;
234 235
}

L
Linus Torvalds 已提交
236 237 238 239 240
static void sysfs_d_iput(struct dentry * dentry, struct inode * inode)
{
	struct sysfs_dirent * sd = dentry->d_fsdata;

	if (sd) {
T
Tejun Heo 已提交
241 242
		/* sd->s_dentry is protected with sysfs_assoc_lock.
		 * This allows sysfs_drop_dentry() to dereference it.
243
		 */
T
Tejun Heo 已提交
244
		spin_lock(&sysfs_assoc_lock);
245 246 247 248 249 250 251 252

		/* The dentry might have been deleted or another
		 * lookup could have happened updating sd->s_dentry to
		 * point the new dentry.  Ignore if it isn't pointing
		 * to this dentry.
		 */
		if (sd->s_dentry == dentry)
			sd->s_dentry = NULL;
T
Tejun Heo 已提交
253
		spin_unlock(&sysfs_assoc_lock);
L
Linus Torvalds 已提交
254 255 256 257 258 259 260 261 262
		sysfs_put(sd);
	}
	iput(inode);
}

static struct dentry_operations sysfs_dentry_ops = {
	.d_iput		= sysfs_d_iput,
};

263
struct sysfs_dirent *sysfs_new_dirent(const char *name, umode_t mode, int type)
L
Linus Torvalds 已提交
264
{
T
Tejun Heo 已提交
265 266 267 268 269 270 271 272
	char *dup_name = NULL;
	struct sysfs_dirent *sd = NULL;

	if (type & SYSFS_COPY_NAME) {
		name = dup_name = kstrdup(name, GFP_KERNEL);
		if (!name)
			goto err_out;
	}
L
Linus Torvalds 已提交
273

274
	sd = kmem_cache_zalloc(sysfs_dir_cachep, GFP_KERNEL);
L
Linus Torvalds 已提交
275
	if (!sd)
T
Tejun Heo 已提交
276
		goto err_out;
L
Linus Torvalds 已提交
277

T
Tejun Heo 已提交
278 279
	if (sysfs_alloc_ino(&sd->s_ino))
		goto err_out;
280

L
Linus Torvalds 已提交
281
	atomic_set(&sd->s_count, 1);
282
	atomic_set(&sd->s_active, 0);
283
	atomic_set(&sd->s_event, 1);
284

T
Tejun Heo 已提交
285
	sd->s_name = name;
286
	sd->s_mode = mode;
287
	sd->s_flags = type;
L
Linus Torvalds 已提交
288 289

	return sd;
T
Tejun Heo 已提交
290 291 292 293 294

 err_out:
	kfree(dup_name);
	kmem_cache_free(sysfs_dir_cachep, sd);
	return NULL;
L
Linus Torvalds 已提交
295 296
}

297 298 299 300 301 302 303 304 305 306 307
/**
 *	sysfs_attach_dentry - associate sysfs_dirent with dentry
 *	@sd: target sysfs_dirent
 *	@dentry: dentry to associate
 *
 *	Associate @sd with @dentry.  This is protected by
 *	sysfs_assoc_lock to avoid race with sysfs_d_iput().
 *
 *	LOCKING:
 *	mutex_lock(sysfs_mutex)
 */
308 309 310 311 312 313
static void sysfs_attach_dentry(struct sysfs_dirent *sd, struct dentry *dentry)
{
	dentry->d_op = &sysfs_dentry_ops;
	dentry->d_fsdata = sysfs_get(sd);

	/* protect sd->s_dentry against sysfs_d_iput */
T
Tejun Heo 已提交
314
	spin_lock(&sysfs_assoc_lock);
315
	sd->s_dentry = dentry;
T
Tejun Heo 已提交
316
	spin_unlock(&sysfs_assoc_lock);
317 318 319 320

	d_rehash(dentry);
}

321 322 323 324 325 326 327 328 329 330 331
/**
 *	sysfs_attach_dirent - attach sysfs_dirent to its parent and dentry
 *	@sd: sysfs_dirent to attach
 *	@parent_sd: parent to attach to (optional)
 *	@dentry: dentry to be associated to @sd (optional)
 *
 *	Attach @sd to @parent_sd and/or @dentry.  Both are optional.
 *
 *	LOCKING:
 *	mutex_lock(sysfs_mutex)
 */
332 333
void sysfs_attach_dirent(struct sysfs_dirent *sd,
			 struct sysfs_dirent *parent_sd, struct dentry *dentry)
334
{
335 336
	if (dentry)
		sysfs_attach_dentry(sd, dentry);
337

T
Tejun Heo 已提交
338 339
	if (parent_sd) {
		sd->s_parent = sysfs_get(parent_sd);
340
		sysfs_link_sibling(sd);
T
Tejun Heo 已提交
341
	}
342 343
}

344 345 346 347 348 349
/**
 *	sysfs_find_dirent - find sysfs_dirent with the given name
 *	@parent_sd: sysfs_dirent to search under
 *	@name: name to look for
 *
 *	Look for sysfs_dirent with name @name under @parent_sd.
350
 *
351
 *	LOCKING:
352
 *	mutex_lock(sysfs_mutex)
353
 *
354 355
 *	RETURNS:
 *	Pointer to sysfs_dirent if found, NULL if not.
356
 */
357 358
struct sysfs_dirent *sysfs_find_dirent(struct sysfs_dirent *parent_sd,
				       const unsigned char *name)
359
{
360 361 362 363 364 365 366
	struct sysfs_dirent *sd;

	for (sd = parent_sd->s_children; sd; sd = sd->s_sibling)
		if (sysfs_type(sd) && !strcmp(sd->s_name, name))
			return sd;
	return NULL;
}
367

368 369 370 371 372 373 374 375 376
/**
 *	sysfs_get_dirent - find and get sysfs_dirent with the given name
 *	@parent_sd: sysfs_dirent to search under
 *	@name: name to look for
 *
 *	Look for sysfs_dirent with name @name under @parent_sd and get
 *	it if found.
 *
 *	LOCKING:
377
 *	Kernel thread context (may sleep).  Grabs sysfs_mutex.
378 379 380 381 382 383 384 385 386
 *
 *	RETURNS:
 *	Pointer to sysfs_dirent if found, NULL if not.
 */
struct sysfs_dirent *sysfs_get_dirent(struct sysfs_dirent *parent_sd,
				      const unsigned char *name)
{
	struct sysfs_dirent *sd;

387
	mutex_lock(&sysfs_mutex);
388 389
	sd = sysfs_find_dirent(parent_sd, name);
	sysfs_get(sd);
390
	mutex_unlock(&sysfs_mutex);
391 392

	return sd;
393 394
}

395 396
static int create_dir(struct kobject *kobj, struct sysfs_dirent *parent_sd,
		      const char *name, struct sysfs_dirent **p_sd)
L
Linus Torvalds 已提交
397
{
398
	struct dentry *parent = parent_sd->s_dentry;
L
Linus Torvalds 已提交
399 400
	int error;
	umode_t mode = S_IFDIR| S_IRWXU | S_IRUGO | S_IXUGO;
401
	struct dentry *dentry;
402
	struct inode *inode;
403
	struct sysfs_dirent *sd;
L
Linus Torvalds 已提交
404

405 406
	mutex_lock(&parent->d_inode->i_mutex);

407
	/* allocate */
408 409 410 411 412 413 414
	dentry = lookup_one_len(name, parent, strlen(name));
	if (IS_ERR(dentry)) {
		error = PTR_ERR(dentry);
		goto out_unlock;
	}

	error = -EEXIST;
415
	if (dentry->d_inode)
416 417
		goto out_dput;

418
	error = -ENOMEM;
419
	sd = sysfs_new_dirent(name, mode, SYSFS_DIR);
420
	if (!sd)
421
		goto out_drop;
422
	sd->s_elem.dir.kobj = kobj;
423

424
	inode = sysfs_get_inode(sd);
425
	if (!inode)
426 427
		goto out_sput;

428 429 430 431 432 433
	if (inode->i_state & I_NEW) {
		inode->i_op = &sysfs_dir_inode_operations;
		inode->i_fop = &sysfs_dir_operations;
		/* directory inodes start off with i_nlink == 2 (for ".") */
		inc_nlink(inode);
	}
434 435

	/* link in */
436 437
	mutex_lock(&sysfs_mutex);

438
	error = -EEXIST;
439 440
	if (sysfs_find_dirent(parent_sd, name)) {
		mutex_unlock(&sysfs_mutex);
441
		goto out_iput;
442
	}
443 444

	sysfs_instantiate(dentry, inode);
445
	inc_nlink(parent->d_inode);
446
	sysfs_attach_dirent(sd, parent_sd, dentry);
447

448 449
	mutex_unlock(&sysfs_mutex);

450
	*p_sd = sd;
451
	error = 0;
452
	goto out_unlock;	/* pin directory dentry in core */
453

454 455
 out_iput:
	iput(inode);
456 457 458 459 460 461 462 463
 out_sput:
	sysfs_put(sd);
 out_drop:
	d_drop(dentry);
 out_dput:
	dput(dentry);
 out_unlock:
	mutex_unlock(&parent->d_inode->i_mutex);
L
Linus Torvalds 已提交
464 465 466
	return error;
}

467 468
int sysfs_create_subdir(struct kobject *kobj, const char *name,
			struct sysfs_dirent **p_sd)
L
Linus Torvalds 已提交
469
{
470
	return create_dir(kobj, kobj->sd, name, p_sd);
L
Linus Torvalds 已提交
471 472 473 474 475
}

/**
 *	sysfs_create_dir - create a directory for an object.
 *	@kobj:		object we're creating directory for. 
476
 *	@shadow_parent:	parent object.
L
Linus Torvalds 已提交
477
 */
478 479
int sysfs_create_dir(struct kobject *kobj,
		     struct sysfs_dirent *shadow_parent_sd)
L
Linus Torvalds 已提交
480
{
481
	struct sysfs_dirent *parent_sd, *sd;
L
Linus Torvalds 已提交
482 483 484 485
	int error = 0;

	BUG_ON(!kobj);

486 487
	if (shadow_parent_sd)
		parent_sd = shadow_parent_sd;
488
	else if (kobj->parent)
489
		parent_sd = kobj->parent->sd;
L
Linus Torvalds 已提交
490
	else if (sysfs_mount && sysfs_mount->mnt_sb)
491
		parent_sd = sysfs_mount->mnt_sb->s_root->d_fsdata;
L
Linus Torvalds 已提交
492 493 494
	else
		return -EFAULT;

495
	error = create_dir(kobj, parent_sd, kobject_name(kobj), &sd);
L
Linus Torvalds 已提交
496
	if (!error)
497
		kobj->sd = sd;
L
Linus Torvalds 已提交
498 499 500 501 502 503 504 505
	return error;
}

static struct dentry * sysfs_lookup(struct inode *dir, struct dentry *dentry,
				struct nameidata *nd)
{
	struct sysfs_dirent * parent_sd = dentry->d_parent->d_fsdata;
	struct sysfs_dirent * sd;
506
	struct bin_attribute *bin_attr;
507 508
	struct inode *inode;
	int found = 0;
L
Linus Torvalds 已提交
509

510
	for (sd = parent_sd->s_children; sd; sd = sd->s_sibling) {
511
		if ((sysfs_type(sd) & SYSFS_NOT_PINNED) &&
512 513
		    !strcmp(sd->s_name, dentry->d_name.name)) {
			found = 1;
L
Linus Torvalds 已提交
514 515 516 517
			break;
		}
	}

518 519 520 521 522
	/* no such entry */
	if (!found)
		return NULL;

	/* attach dentry and inode */
523
	inode = sysfs_get_inode(sd);
524 525 526
	if (!inode)
		return ERR_PTR(-ENOMEM);

527 528
	mutex_lock(&sysfs_mutex);

529 530
	if (inode->i_state & I_NEW) {
		/* initialize inode according to type */
531 532
		switch (sysfs_type(sd)) {
		case SYSFS_KOBJ_ATTR:
533 534
			inode->i_size = PAGE_SIZE;
			inode->i_fop = &sysfs_file_operations;
535 536 537
			break;
		case SYSFS_KOBJ_BIN_ATTR:
			bin_attr = sd->s_elem.bin_attr.bin_attr;
538 539
			inode->i_size = bin_attr->size;
			inode->i_fop = &bin_fops;
540 541
			break;
		case SYSFS_KOBJ_LINK:
542
			inode->i_op = &sysfs_symlink_inode_operations;
543 544 545 546
			break;
		default:
			BUG();
		}
547
	}
548 549 550 551

	sysfs_instantiate(dentry, inode);
	sysfs_attach_dentry(sd, dentry);

552 553
	mutex_unlock(&sysfs_mutex);

554
	return NULL;
L
Linus Torvalds 已提交
555 556
}

557
const struct inode_operations sysfs_dir_inode_operations = {
L
Linus Torvalds 已提交
558
	.lookup		= sysfs_lookup,
559
	.setattr	= sysfs_setattr,
L
Linus Torvalds 已提交
560 561
};

562
static void remove_dir(struct sysfs_dirent *sd)
L
Linus Torvalds 已提交
563
{
564
	mutex_lock(&sysfs_mutex);
565
	sysfs_unlink_sibling(sd);
566
	sd->s_flags |= SYSFS_FLAG_REMOVED;
567
	mutex_unlock(&sysfs_mutex);
L
Linus Torvalds 已提交
568

569
	pr_debug(" o %s removing done\n", sd->s_name);
L
Linus Torvalds 已提交
570

571
	sysfs_drop_dentry(sd);
572 573
	sysfs_deactivate(sd);
	sysfs_put(sd);
L
Linus Torvalds 已提交
574 575
}

576
void sysfs_remove_subdir(struct sysfs_dirent *sd)
L
Linus Torvalds 已提交
577
{
578
	remove_dir(sd);
L
Linus Torvalds 已提交
579 580 581
}


582
static void __sysfs_remove_dir(struct sysfs_dirent *dir_sd)
L
Linus Torvalds 已提交
583
{
584 585
	struct sysfs_dirent *removed = NULL;
	struct sysfs_dirent **pos;
L
Linus Torvalds 已提交
586

587
	if (!dir_sd)
L
Linus Torvalds 已提交
588 589
		return;

590
	pr_debug("sysfs %s: removing dir\n", dir_sd->s_name);
591
	mutex_lock(&sysfs_mutex);
592
	pos = &dir_sd->s_children;
593 594 595
	while (*pos) {
		struct sysfs_dirent *sd = *pos;

596
		if (sysfs_type(sd) && (sysfs_type(sd) & SYSFS_NOT_PINNED)) {
597
			sd->s_flags |= SYSFS_FLAG_REMOVED;
598 599 600 601 602
			*pos = sd->s_sibling;
			sd->s_sibling = removed;
			removed = sd;
		} else
			pos = &(*pos)->s_sibling;
L
Linus Torvalds 已提交
603
	}
604
	mutex_unlock(&sysfs_mutex);
L
Linus Torvalds 已提交
605

606 607 608 609 610 611
	while (removed) {
		struct sysfs_dirent *sd = removed;

		removed = sd->s_sibling;
		sd->s_sibling = NULL;

612
		sysfs_drop_dentry(sd);
613 614 615 616
		sysfs_deactivate(sd);
		sysfs_put(sd);
	}

617
	remove_dir(dir_sd);
618 619 620 621 622 623 624 625 626 627 628 629 630
}

/**
 *	sysfs_remove_dir - remove an object's directory.
 *	@kobj:	object.
 *
 *	The only thing special about this is that we remove any files in
 *	the directory before we remove the directory, and we've inlined
 *	what used to be sysfs_rmdir() below, instead of calling separately.
 */

void sysfs_remove_dir(struct kobject * kobj)
{
631
	struct sysfs_dirent *sd = kobj->sd;
632

T
Tejun Heo 已提交
633
	spin_lock(&sysfs_assoc_lock);
634
	kobj->sd = NULL;
T
Tejun Heo 已提交
635
	spin_unlock(&sysfs_assoc_lock);
636

637
	__sysfs_remove_dir(sd);
L
Linus Torvalds 已提交
638 639
}

640
int sysfs_rename_dir(struct kobject *kobj, struct sysfs_dirent *new_parent_sd,
641
		     const char *new_name)
L
Linus Torvalds 已提交
642
{
643 644
	struct sysfs_dirent *sd = kobj->sd;
	struct dentry *new_parent = new_parent_sd->s_dentry;
T
Tejun Heo 已提交
645 646
	struct dentry *new_dentry;
	char *dup_name;
647
	int error;
L
Linus Torvalds 已提交
648

649
	if (!new_parent_sd)
650
		return -EFAULT;
L
Linus Torvalds 已提交
651

652
	mutex_lock(&new_parent->d_inode->i_mutex);
L
Linus Torvalds 已提交
653

654
	new_dentry = lookup_one_len(new_name, new_parent, strlen(new_name));
655 656 657
	if (IS_ERR(new_dentry)) {
		error = PTR_ERR(new_dentry);
		goto out_unlock;
L
Linus Torvalds 已提交
658
	}
659 660 661 662 663 664

	/* By allowing two different directories with the same
	 * d_parent we allow this routine to move between different
	 * shadows of the same directory
	 */
	error = -EINVAL;
665
	if (sd->s_parent->s_dentry->d_inode != new_parent->d_inode ||
666
	    new_dentry->d_parent->d_inode != new_parent->d_inode ||
667
	    new_dentry == sd->s_dentry)
668 669 670 671 672 673
		goto out_dput;

	error = -EEXIST;
	if (new_dentry->d_inode)
		goto out_dput;

T
Tejun Heo 已提交
674 675 676 677 678 679
	/* rename kobject and sysfs_dirent */
	error = -ENOMEM;
	new_name = dup_name = kstrdup(new_name, GFP_KERNEL);
	if (!new_name)
		goto out_drop;

680 681
	error = kobject_set_name(kobj, "%s", new_name);
	if (error)
T
Tejun Heo 已提交
682
		goto out_free;
683

T
Tejun Heo 已提交
684 685 686 687
	kfree(sd->s_name);
	sd->s_name = new_name;

	/* move under the new parent */
688
	d_add(new_dentry, NULL);
689
	d_move(sd->s_dentry, new_dentry);
690

691 692
	mutex_lock(&sysfs_mutex);

693
	sysfs_unlink_sibling(sd);
694
	sysfs_get(new_parent_sd);
695
	sysfs_put(sd->s_parent);
696
	sd->s_parent = new_parent_sd;
697
	sysfs_link_sibling(sd);
698

699 700
	mutex_unlock(&sysfs_mutex);

701 702 703
	error = 0;
	goto out_unlock;

T
Tejun Heo 已提交
704 705
 out_free:
	kfree(dup_name);
706 707 708 709 710
 out_drop:
	d_drop(new_dentry);
 out_dput:
	dput(new_dentry);
 out_unlock:
711
	mutex_unlock(&new_parent->d_inode->i_mutex);
L
Linus Torvalds 已提交
712 713 714
	return error;
}

715 716 717 718 719 720 721
int sysfs_move_dir(struct kobject *kobj, struct kobject *new_parent)
{
	struct dentry *old_parent_dentry, *new_parent_dentry, *new_dentry;
	struct sysfs_dirent *new_parent_sd, *sd;
	int error;

	old_parent_dentry = kobj->parent ?
722
		kobj->parent->sd->s_dentry : sysfs_mount->mnt_sb->s_root;
723
	new_parent_dentry = new_parent ?
724
		new_parent->sd->s_dentry : sysfs_mount->mnt_sb->s_root;
725

M
Mark Lord 已提交
726 727
	if (old_parent_dentry->d_inode == new_parent_dentry->d_inode)
		return 0;	/* nothing to move */
728 729 730 731 732 733 734 735
again:
	mutex_lock(&old_parent_dentry->d_inode->i_mutex);
	if (!mutex_trylock(&new_parent_dentry->d_inode->i_mutex)) {
		mutex_unlock(&old_parent_dentry->d_inode->i_mutex);
		goto again;
	}

	new_parent_sd = new_parent_dentry->d_fsdata;
736
	sd = kobj->sd;
737 738 739 740 741 742 743 744 745

	new_dentry = lookup_one_len(kobj->name, new_parent_dentry,
				    strlen(kobj->name));
	if (IS_ERR(new_dentry)) {
		error = PTR_ERR(new_dentry);
		goto out;
	} else
		error = 0;
	d_add(new_dentry, NULL);
746
	d_move(sd->s_dentry, new_dentry);
747 748 749
	dput(new_dentry);

	/* Remove from old parent's list and insert into new parent's list. */
750 751
	mutex_lock(&sysfs_mutex);

752
	sysfs_unlink_sibling(sd);
753 754 755
	sysfs_get(new_parent_sd);
	sysfs_put(sd->s_parent);
	sd->s_parent = new_parent_sd;
756
	sysfs_link_sibling(sd);
757

758
	mutex_unlock(&sysfs_mutex);
759 760 761 762 763 764 765
out:
	mutex_unlock(&new_parent_dentry->d_inode->i_mutex);
	mutex_unlock(&old_parent_dentry->d_inode->i_mutex);

	return error;
}

L
Linus Torvalds 已提交
766 767
static int sysfs_dir_open(struct inode *inode, struct file *file)
{
768
	struct dentry * dentry = file->f_path.dentry;
L
Linus Torvalds 已提交
769
	struct sysfs_dirent * parent_sd = dentry->d_fsdata;
770
	struct sysfs_dirent * sd;
L
Linus Torvalds 已提交
771

772
	sd = sysfs_new_dirent("_DIR_", 0, 0);
773 774
	if (sd) {
		mutex_lock(&sysfs_mutex);
775
		sysfs_attach_dirent(sd, parent_sd, NULL);
776 777
		mutex_unlock(&sysfs_mutex);
	}
L
Linus Torvalds 已提交
778

779 780
	file->private_data = sd;
	return sd ? 0 : -ENOMEM;
L
Linus Torvalds 已提交
781 782 783 784 785 786
}

static int sysfs_dir_close(struct inode *inode, struct file *file)
{
	struct sysfs_dirent * cursor = file->private_data;

787
	mutex_lock(&sysfs_mutex);
788
	sysfs_unlink_sibling(cursor);
789
	mutex_unlock(&sysfs_mutex);
L
Linus Torvalds 已提交
790 791 792 793 794 795 796 797 798 799 800 801 802 803

	release_sysfs_dirent(cursor);

	return 0;
}

/* Relationship between s_mode and the DT_xxx types */
static inline unsigned char dt_type(struct sysfs_dirent *sd)
{
	return (sd->s_mode >> 12) & 15;
}

static int sysfs_readdir(struct file * filp, void * dirent, filldir_t filldir)
{
804
	struct dentry *dentry = filp->f_path.dentry;
L
Linus Torvalds 已提交
805 806
	struct sysfs_dirent * parent_sd = dentry->d_fsdata;
	struct sysfs_dirent *cursor = filp->private_data;
807
	struct sysfs_dirent **pos;
L
Linus Torvalds 已提交
808 809 810 811 812
	ino_t ino;
	int i = filp->f_pos;

	switch (i) {
		case 0:
813
			ino = parent_sd->s_ino;
L
Linus Torvalds 已提交
814 815 816 817 818 819
			if (filldir(dirent, ".", 1, i, ino, DT_DIR) < 0)
				break;
			filp->f_pos++;
			i++;
			/* fallthrough */
		case 1:
T
Tejun Heo 已提交
820 821 822 823
			if (parent_sd->s_parent)
				ino = parent_sd->s_parent->s_ino;
			else
				ino = parent_sd->s_ino;
L
Linus Torvalds 已提交
824 825 826 827 828 829
			if (filldir(dirent, "..", 2, i, ino, DT_DIR) < 0)
				break;
			filp->f_pos++;
			i++;
			/* fallthrough */
		default:
830 831
			mutex_lock(&sysfs_mutex);

832 833 834 835 836 837 838
			pos = &parent_sd->s_children;
			while (*pos != cursor)
				pos = &(*pos)->s_sibling;

			/* unlink cursor */
			*pos = cursor->s_sibling;

A
Akinobu Mita 已提交
839
			if (filp->f_pos == 2)
840
				pos = &parent_sd->s_children;
A
Akinobu Mita 已提交
841

842 843
			for ( ; *pos; pos = &(*pos)->s_sibling) {
				struct sysfs_dirent *next = *pos;
L
Linus Torvalds 已提交
844 845 846
				const char * name;
				int len;

847
				if (!sysfs_type(next))
L
Linus Torvalds 已提交
848 849
					continue;

T
Tejun Heo 已提交
850
				name = next->s_name;
L
Linus Torvalds 已提交
851
				len = strlen(name);
852
				ino = next->s_ino;
L
Linus Torvalds 已提交
853 854 855

				if (filldir(dirent, name, len, filp->f_pos, ino,
						 dt_type(next)) < 0)
856
					break;
L
Linus Torvalds 已提交
857 858 859

				filp->f_pos++;
			}
860 861 862 863

			/* put cursor back in */
			cursor->s_sibling = *pos;
			*pos = cursor;
864 865

			mutex_unlock(&sysfs_mutex);
L
Linus Torvalds 已提交
866 867 868 869 870 871
	}
	return 0;
}

static loff_t sysfs_dir_lseek(struct file * file, loff_t offset, int origin)
{
872
	struct dentry * dentry = file->f_path.dentry;
L
Linus Torvalds 已提交
873 874 875 876 877 878 879 880 881 882 883

	switch (origin) {
		case 1:
			offset += file->f_pos;
		case 0:
			if (offset >= 0)
				break;
		default:
			return -EINVAL;
	}
	if (offset != file->f_pos) {
884 885
		mutex_lock(&sysfs_mutex);

L
Linus Torvalds 已提交
886 887 888 889
		file->f_pos = offset;
		if (file->f_pos >= 2) {
			struct sysfs_dirent *sd = dentry->d_fsdata;
			struct sysfs_dirent *cursor = file->private_data;
890
			struct sysfs_dirent **pos;
L
Linus Torvalds 已提交
891 892
			loff_t n = file->f_pos - 2;

893 894 895 896 897
			sysfs_unlink_sibling(cursor);

			pos = &sd->s_children;
			while (n && *pos) {
				struct sysfs_dirent *next = *pos;
898
				if (sysfs_type(next))
L
Linus Torvalds 已提交
899
					n--;
900
				pos = &(*pos)->s_sibling;
L
Linus Torvalds 已提交
901
			}
902 903 904

			cursor->s_sibling = *pos;
			*pos = cursor;
L
Linus Torvalds 已提交
905
		}
906 907

		mutex_unlock(&sysfs_mutex);
L
Linus Torvalds 已提交
908
	}
909

L
Linus Torvalds 已提交
910 911 912
	return offset;
}

913 914 915 916 917 918 919 920 921 922 923 924

/**
 *	sysfs_make_shadowed_dir - Setup so a directory can be shadowed
 *	@kobj:	object we're creating shadow of.
 */

int sysfs_make_shadowed_dir(struct kobject *kobj,
	void * (*follow_link)(struct dentry *, struct nameidata *))
{
	struct inode *inode;
	struct inode_operations *i_op;

925
	inode = kobj->sd->s_dentry->d_inode;
926 927 928 929 930 931 932 933 934 935 936 937 938 939 940 941 942 943 944 945 946 947 948 949 950 951
	if (inode->i_op != &sysfs_dir_inode_operations)
		return -EINVAL;

	i_op = kmalloc(sizeof(*i_op), GFP_KERNEL);
	if (!i_op)
		return -ENOMEM;

	memcpy(i_op, &sysfs_dir_inode_operations, sizeof(*i_op));
	i_op->follow_link = follow_link;

	/* Locking of inode->i_op?
	 * Since setting i_op is a single word write and they
	 * are atomic we should be ok here.
	 */
	inode->i_op = i_op;
	return 0;
}

/**
 *	sysfs_create_shadow_dir - create a shadow directory for an object.
 *	@kobj:	object we're creating directory for.
 *
 *	sysfs_make_shadowed_dir must already have been called on this
 *	directory.
 */

952
struct sysfs_dirent *sysfs_create_shadow_dir(struct kobject *kobj)
953
{
954
	struct dentry *dir = kobj->sd->s_dentry;
T
Tejun Heo 已提交
955 956 957 958
	struct inode *inode = dir->d_inode;
	struct dentry *parent = dir->d_parent;
	struct sysfs_dirent *parent_sd = parent->d_fsdata;
	struct dentry *shadow;
959 960
	struct sysfs_dirent *sd;

961
	sd = ERR_PTR(-EINVAL);
962 963 964 965 966 967 968
	if (!sysfs_is_shadowed_inode(inode))
		goto out;

	shadow = d_alloc(parent, &dir->d_name);
	if (!shadow)
		goto nomem;

969
	sd = sysfs_new_dirent("_SHADOW_", inode->i_mode, SYSFS_DIR);
970 971
	if (!sd)
		goto nomem;
972
	sd->s_elem.dir.kobj = kobj;
T
Tejun Heo 已提交
973 974
	/* point to parent_sd but don't attach to it */
	sd->s_parent = sysfs_get(parent_sd);
975
	mutex_lock(&sysfs_mutex);
976
	sysfs_attach_dirent(sd, NULL, shadow);
977
	mutex_unlock(&sysfs_mutex);
978 979 980 981 982 983 984 985

	d_instantiate(shadow, igrab(inode));
	inc_nlink(inode);
	inc_nlink(parent->d_inode);

	dget(shadow);		/* Extra count - pin the dentry in core */

out:
986
	return sd;
987 988
nomem:
	dput(shadow);
989
	sd = ERR_PTR(-ENOMEM);
990 991 992 993 994
	goto out;
}

/**
 *	sysfs_remove_shadow_dir - remove an object's directory.
995
 *	@shadow_sd: sysfs_dirent of shadow directory
996 997 998 999 1000 1001
 *
 *	The only thing special about this is that we remove any files in
 *	the directory before we remove the directory, and we've inlined
 *	what used to be sysfs_rmdir() below, instead of calling separately.
 */

1002
void sysfs_remove_shadow_dir(struct sysfs_dirent *shadow_sd)
1003
{
1004
	__sysfs_remove_dir(shadow_sd);
1005 1006
}

1007
const struct file_operations sysfs_dir_operations = {
L
Linus Torvalds 已提交
1008 1009 1010 1011 1012 1013
	.open		= sysfs_dir_open,
	.release	= sysfs_dir_close,
	.llseek		= sysfs_dir_lseek,
	.read		= generic_read_dir,
	.readdir	= sysfs_readdir,
};