sched_debug.c 11.9 KB
Newer Older
I
Ingo Molnar 已提交
1 2 3 4 5 6 7 8 9 10 11 12 13 14 15 16 17 18 19 20 21 22 23 24 25 26 27 28 29 30
/*
 * kernel/time/sched_debug.c
 *
 * Print the CFS rbtree
 *
 * Copyright(C) 2007, Red Hat, Inc., Ingo Molnar
 *
 * This program is free software; you can redistribute it and/or modify
 * it under the terms of the GNU General Public License version 2 as
 * published by the Free Software Foundation.
 */

#include <linux/proc_fs.h>
#include <linux/sched.h>
#include <linux/seq_file.h>
#include <linux/kallsyms.h>
#include <linux/utsname.h>

/*
 * This allows printing both to /proc/sched_debug and
 * to the console
 */
#define SEQ_printf(m, x...)			\
 do {						\
	if (m)					\
		seq_printf(m, x);		\
	else					\
		printk(x);			\
 } while (0)

I
Ingo Molnar 已提交
31 32 33
/*
 * Ease the printing of nsec fields:
 */
I
Ingo Molnar 已提交
34
static long long nsec_high(unsigned long long nsec)
I
Ingo Molnar 已提交
35
{
I
Ingo Molnar 已提交
36
	if ((long long)nsec < 0) {
I
Ingo Molnar 已提交
37 38 39 40 41 42 43 44 45
		nsec = -nsec;
		do_div(nsec, 1000000);
		return -nsec;
	}
	do_div(nsec, 1000000);

	return nsec;
}

I
Ingo Molnar 已提交
46
static unsigned long nsec_low(unsigned long long nsec)
I
Ingo Molnar 已提交
47
{
I
Ingo Molnar 已提交
48
	if ((long long)nsec < 0)
I
Ingo Molnar 已提交
49 50 51 52 53 54 55
		nsec = -nsec;

	return do_div(nsec, 1000000);
}

#define SPLIT_NS(x) nsec_high(x), nsec_low(x)

56 57 58 59 60 61 62 63 64 65 66 67 68 69 70 71 72 73 74 75 76 77 78 79 80 81 82 83 84 85 86 87 88 89
#ifdef CONFIG_FAIR_GROUP_SCHED
static void print_cfs_group_stats(struct seq_file *m, int cpu,
		struct task_group *tg)
{
	struct sched_entity *se = tg->se[cpu];
	if (!se)
		return;

#define P(F) \
	SEQ_printf(m, "  .%-30s: %lld\n", #F, (long long)F)
#define PN(F) \
	SEQ_printf(m, "  .%-30s: %lld.%06ld\n", #F, SPLIT_NS((long long)F))

	PN(se->exec_start);
	PN(se->vruntime);
	PN(se->sum_exec_runtime);
#ifdef CONFIG_SCHEDSTATS
	PN(se->wait_start);
	PN(se->sleep_start);
	PN(se->block_start);
	PN(se->sleep_max);
	PN(se->block_max);
	PN(se->exec_max);
	PN(se->slice_max);
	PN(se->wait_max);
	PN(se->wait_sum);
	P(se->wait_count);
#endif
	P(se->load.weight);
#undef PN
#undef P
}
#endif

I
Ingo Molnar 已提交
90
static void
91
print_task(struct seq_file *m, struct rq *rq, struct task_struct *p)
I
Ingo Molnar 已提交
92 93 94 95 96 97
{
	if (rq->curr == p)
		SEQ_printf(m, "R");
	else
		SEQ_printf(m, " ");

I
Ingo Molnar 已提交
98
	SEQ_printf(m, "%15s %5d %9Ld.%06ld %9Ld %5d ",
I
Ingo Molnar 已提交
99
		p->comm, p->pid,
I
Ingo Molnar 已提交
100
		SPLIT_NS(p->se.vruntime),
I
Ingo Molnar 已提交
101
		(long long)(p->nvcsw + p->nivcsw),
102
		p->prio);
I
Ingo Molnar 已提交
103
#ifdef CONFIG_SCHEDSTATS
104
	SEQ_printf(m, "%9Ld.%06ld %9Ld.%06ld %9Ld.%06ld",
I
Ingo Molnar 已提交
105 106 107
		SPLIT_NS(p->se.vruntime),
		SPLIT_NS(p->se.sum_exec_runtime),
		SPLIT_NS(p->se.sum_sleep_runtime));
I
Ingo Molnar 已提交
108
#else
109
	SEQ_printf(m, "%15Ld %15Ld %15Ld.%06ld %15Ld.%06ld %15Ld.%06ld",
I
Ingo Molnar 已提交
110
		0LL, 0LL, 0LL, 0L, 0LL, 0L, 0LL, 0L);
I
Ingo Molnar 已提交
111
#endif
112 113 114 115 116 117 118 119 120 121

#ifdef CONFIG_CGROUP_SCHED
	{
		char path[64];

		cgroup_path(task_group(p)->css.cgroup, path, sizeof(path));
		SEQ_printf(m, " %s", path);
	}
#endif
	SEQ_printf(m, "\n");
I
Ingo Molnar 已提交
122 123
}

124
static void print_rq(struct seq_file *m, struct rq *rq, int rq_cpu)
I
Ingo Molnar 已提交
125 126
{
	struct task_struct *g, *p;
P
Peter Zijlstra 已提交
127
	unsigned long flags;
I
Ingo Molnar 已提交
128 129 130

	SEQ_printf(m,
	"\nrunnable tasks:\n"
131 132
	"            task   PID         tree-key  switches  prio"
	"     exec-runtime         sum-exec        sum-sleep\n"
133
	"------------------------------------------------------"
134
	"----------------------------------------------------\n");
I
Ingo Molnar 已提交
135

P
Peter Zijlstra 已提交
136
	read_lock_irqsave(&tasklist_lock, flags);
I
Ingo Molnar 已提交
137 138 139 140 141

	do_each_thread(g, p) {
		if (!p->se.on_rq || task_cpu(p) != rq_cpu)
			continue;

142
		print_task(m, rq, p);
I
Ingo Molnar 已提交
143 144
	} while_each_thread(g, p);

P
Peter Zijlstra 已提交
145
	read_unlock_irqrestore(&tasklist_lock, flags);
I
Ingo Molnar 已提交
146 147
}

148 149 150 151 152 153 154 155 156 157 158 159 160
#if defined(CONFIG_CGROUP_SCHED) && \
	(defined(CONFIG_FAIR_GROUP_SCHED) || defined(CONFIG_RT_GROUP_SCHED))
static void task_group_path(struct task_group *tg, char *buf, int buflen)
{
	/* may be NULL if the underlying cgroup isn't fully-created yet */
	if (!tg->css.cgroup) {
		buf[0] = '\0';
		return;
	}
	cgroup_path(tg->css.cgroup, buf, buflen);
}
#endif

161
void print_cfs_rq(struct seq_file *m, int cpu, struct cfs_rq *cfs_rq)
I
Ingo Molnar 已提交
162
{
I
Ingo Molnar 已提交
163 164
	s64 MIN_vruntime = -1, min_vruntime, max_vruntime = -1,
		spread, rq0_min_vruntime, spread0;
165
	struct rq *rq = cpu_rq(cpu);
I
Ingo Molnar 已提交
166 167 168
	struct sched_entity *last;
	unsigned long flags;

169
#if defined(CONFIG_CGROUP_SCHED) && defined(CONFIG_FAIR_GROUP_SCHED)
170
	char path[128];
171 172
	struct task_group *tg = cfs_rq->tg;

173
	task_group_path(tg, path, sizeof(path));
174 175

	SEQ_printf(m, "\ncfs_rq[%d]:%s\n", cpu, path);
176 177 178 179 180
#elif defined(CONFIG_USER_SCHED) && defined(CONFIG_FAIR_GROUP_SCHED)
	{
		uid_t uid = cfs_rq->tg->uid;
		SEQ_printf(m, "\ncfs_rq[%d] for UID: %u\n", cpu, uid);
	}
181 182
#else
	SEQ_printf(m, "\ncfs_rq[%d]:\n", cpu);
183
#endif
I
Ingo Molnar 已提交
184 185
	SEQ_printf(m, "  .%-30s: %Ld.%06ld\n", "exec_clock",
			SPLIT_NS(cfs_rq->exec_clock));
I
Ingo Molnar 已提交
186 187 188 189 190 191 192

	spin_lock_irqsave(&rq->lock, flags);
	if (cfs_rq->rb_leftmost)
		MIN_vruntime = (__pick_next_entity(cfs_rq))->vruntime;
	last = __pick_last_entity(cfs_rq);
	if (last)
		max_vruntime = last->vruntime;
P
Peter Zijlstra 已提交
193
	min_vruntime = cfs_rq->min_vruntime;
194
	rq0_min_vruntime = cpu_rq(0)->cfs.min_vruntime;
I
Ingo Molnar 已提交
195
	spin_unlock_irqrestore(&rq->lock, flags);
I
Ingo Molnar 已提交
196 197 198 199 200 201
	SEQ_printf(m, "  .%-30s: %Ld.%06ld\n", "MIN_vruntime",
			SPLIT_NS(MIN_vruntime));
	SEQ_printf(m, "  .%-30s: %Ld.%06ld\n", "min_vruntime",
			SPLIT_NS(min_vruntime));
	SEQ_printf(m, "  .%-30s: %Ld.%06ld\n", "max_vruntime",
			SPLIT_NS(max_vruntime));
I
Ingo Molnar 已提交
202
	spread = max_vruntime - MIN_vruntime;
I
Ingo Molnar 已提交
203 204
	SEQ_printf(m, "  .%-30s: %Ld.%06ld\n", "spread",
			SPLIT_NS(spread));
I
Ingo Molnar 已提交
205
	spread0 = min_vruntime - rq0_min_vruntime;
I
Ingo Molnar 已提交
206 207
	SEQ_printf(m, "  .%-30s: %Ld.%06ld\n", "spread0",
			SPLIT_NS(spread0));
208 209
	SEQ_printf(m, "  .%-30s: %ld\n", "nr_running", cfs_rq->nr_running);
	SEQ_printf(m, "  .%-30s: %ld\n", "load", cfs_rq->load.weight);
210

P
Peter Zijlstra 已提交
211
	SEQ_printf(m, "  .%-30s: %d\n", "nr_spread_over",
P
Peter Zijlstra 已提交
212
			cfs_rq->nr_spread_over);
213 214 215 216
#ifdef CONFIG_FAIR_GROUP_SCHED
#ifdef CONFIG_SMP
	SEQ_printf(m, "  .%-30s: %lu\n", "shares", cfs_rq->shares);
#endif
217
	print_cfs_group_stats(m, cpu, cfs_rq->tg);
218
#endif
I
Ingo Molnar 已提交
219 220
}

221 222 223
void print_rt_rq(struct seq_file *m, int cpu, struct rt_rq *rt_rq)
{
#if defined(CONFIG_CGROUP_SCHED) && defined(CONFIG_RT_GROUP_SCHED)
224
	char path[128];
225 226
	struct task_group *tg = rt_rq->tg;

227
	task_group_path(tg, path, sizeof(path));
228 229 230 231 232 233 234 235 236 237 238 239 240 241 242 243 244 245 246 247 248

	SEQ_printf(m, "\nrt_rq[%d]:%s\n", cpu, path);
#else
	SEQ_printf(m, "\nrt_rq[%d]:\n", cpu);
#endif


#define P(x) \
	SEQ_printf(m, "  .%-30s: %Ld\n", #x, (long long)(rt_rq->x))
#define PN(x) \
	SEQ_printf(m, "  .%-30s: %Ld.%06ld\n", #x, SPLIT_NS(rt_rq->x))

	P(rt_nr_running);
	P(rt_throttled);
	PN(rt_time);
	PN(rt_runtime);

#undef PN
#undef P
}

249
static void print_cpu(struct seq_file *m, int cpu)
I
Ingo Molnar 已提交
250
{
251
	struct rq *rq = cpu_rq(cpu);
I
Ingo Molnar 已提交
252 253 254 255 256 257 258 259 260 261 262 263 264 265

#ifdef CONFIG_X86
	{
		unsigned int freq = cpu_khz ? : 1;

		SEQ_printf(m, "\ncpu#%d, %u.%03u MHz\n",
			   cpu, freq / 1000, (freq % 1000));
	}
#else
	SEQ_printf(m, "\ncpu#%d\n", cpu);
#endif

#define P(x) \
	SEQ_printf(m, "  .%-30s: %Ld\n", #x, (long long)(rq->x))
I
Ingo Molnar 已提交
266 267
#define PN(x) \
	SEQ_printf(m, "  .%-30s: %Ld.%06ld\n", #x, SPLIT_NS(rq->x))
I
Ingo Molnar 已提交
268 269 270

	P(nr_running);
	SEQ_printf(m, "  .%-30s: %lu\n", "load",
271
		   rq->load.weight);
I
Ingo Molnar 已提交
272 273 274
	P(nr_switches);
	P(nr_load_updates);
	P(nr_uninterruptible);
I
Ingo Molnar 已提交
275
	PN(next_balance);
I
Ingo Molnar 已提交
276
	P(curr->pid);
I
Ingo Molnar 已提交
277
	PN(clock);
I
Ingo Molnar 已提交
278 279 280 281 282 283
	P(cpu_load[0]);
	P(cpu_load[1]);
	P(cpu_load[2]);
	P(cpu_load[3]);
	P(cpu_load[4]);
#undef P
I
Ingo Molnar 已提交
284
#undef PN
I
Ingo Molnar 已提交
285

P
Peter Zijlstra 已提交
286 287
#ifdef CONFIG_SCHEDSTATS
#define P(n) SEQ_printf(m, "  .%-30s: %d\n", #n, rq->n);
M
Mike Galbraith 已提交
288
#define P64(n) SEQ_printf(m, "  .%-30s: %Ld\n", #n, rq->n);
P
Peter Zijlstra 已提交
289 290 291 292 293 294

	P(yld_count);

	P(sched_switch);
	P(sched_count);
	P(sched_goidle);
M
Mike Galbraith 已提交
295 296 297
#ifdef CONFIG_SMP
	P64(avg_idle);
#endif
P
Peter Zijlstra 已提交
298 299 300 301 302 303 304 305

	P(ttwu_count);
	P(ttwu_local);

	P(bkl_count);

#undef P
#endif
306
	print_cfs_stats(m, cpu);
307
	print_rt_stats(m, cpu);
I
Ingo Molnar 已提交
308

309
	print_rq(m, rq, cpu);
I
Ingo Molnar 已提交
310 311
}

312 313 314 315 316 317
static const char *sched_tunable_scaling_names[] = {
	"none",
	"logaritmic",
	"linear"
};

I
Ingo Molnar 已提交
318 319 320 321 322
static int sched_debug_show(struct seq_file *m, void *v)
{
	u64 now = ktime_to_ns(ktime_get());
	int cpu;

323
	SEQ_printf(m, "Sched Debug Version: v0.09, %s %.*s\n",
I
Ingo Molnar 已提交
324 325 326 327
		init_utsname()->release,
		(int)strcspn(init_utsname()->version, " "),
		init_utsname()->version);

I
Ingo Molnar 已提交
328
	SEQ_printf(m, "now at %Lu.%06ld msecs\n", SPLIT_NS(now));
I
Ingo Molnar 已提交
329

I
Ingo Molnar 已提交
330
#define P(x) \
331
	SEQ_printf(m, "  .%-40s: %Ld\n", #x, (long long)(x))
I
Ingo Molnar 已提交
332
#define PN(x) \
333
	SEQ_printf(m, "  .%-40s: %Ld.%06ld\n", #x, SPLIT_NS(x))
334
	P(jiffies);
I
Ingo Molnar 已提交
335
	PN(sysctl_sched_latency);
336
	PN(sysctl_sched_min_granularity);
I
Ingo Molnar 已提交
337 338 339 340 341 342
	PN(sysctl_sched_wakeup_granularity);
	PN(sysctl_sched_child_runs_first);
	P(sysctl_sched_features);
#undef PN
#undef P

343 344 345 346
	SEQ_printf(m, "  .%-40s: %d (%s)\n", "sysctl_sched_tunable_scaling",
		sysctl_sched_tunable_scaling,
		sched_tunable_scaling_names[sysctl_sched_tunable_scaling]);

I
Ingo Molnar 已提交
347
	for_each_online_cpu(cpu)
348
		print_cpu(m, cpu);
I
Ingo Molnar 已提交
349 350 351 352 353 354

	SEQ_printf(m, "\n");

	return 0;
}

355
static void sysrq_sched_debug_show(void)
I
Ingo Molnar 已提交
356 357 358 359 360 361 362 363 364
{
	sched_debug_show(NULL, NULL);
}

static int sched_debug_open(struct inode *inode, struct file *filp)
{
	return single_open(filp, sched_debug_show, NULL);
}

365
static const struct file_operations sched_debug_fops = {
I
Ingo Molnar 已提交
366 367 368
	.open		= sched_debug_open,
	.read		= seq_read,
	.llseek		= seq_lseek,
369
	.release	= single_release,
I
Ingo Molnar 已提交
370 371 372 373 374 375
};

static int __init init_sched_debug_procfs(void)
{
	struct proc_dir_entry *pe;

376
	pe = proc_create("sched_debug", 0444, NULL, &sched_debug_fops);
I
Ingo Molnar 已提交
377 378 379 380 381 382 383 384 385
	if (!pe)
		return -ENOMEM;
	return 0;
}

__initcall(init_sched_debug_procfs);

void proc_sched_show_task(struct task_struct *p, struct seq_file *m)
{
386
	unsigned long nr_switches;
I
Ingo Molnar 已提交
387 388 389 390 391 392 393 394 395
	unsigned long flags;
	int num_threads = 1;

	if (lock_task_sighand(p, &flags)) {
		num_threads = atomic_read(&p->signal->count);
		unlock_task_sighand(p, &flags);
	}

	SEQ_printf(m, "%s (%d, #threads: %d)\n", p->comm, p->pid, num_threads);
396 397
	SEQ_printf(m,
		"---------------------------------------------------------\n");
398 399
#define __P(F) \
	SEQ_printf(m, "%-35s:%21Ld\n", #F, (long long)F)
I
Ingo Molnar 已提交
400
#define P(F) \
401
	SEQ_printf(m, "%-35s:%21Ld\n", #F, (long long)p->F)
402 403
#define __PN(F) \
	SEQ_printf(m, "%-35s:%14Ld.%06ld\n", #F, SPLIT_NS((long long)F))
I
Ingo Molnar 已提交
404
#define PN(F) \
405
	SEQ_printf(m, "%-35s:%14Ld.%06ld\n", #F, SPLIT_NS((long long)p->F))
I
Ingo Molnar 已提交
406

I
Ingo Molnar 已提交
407 408 409
	PN(se.exec_start);
	PN(se.vruntime);
	PN(se.sum_exec_runtime);
I
Ingo Molnar 已提交
410
	PN(se.avg_overlap);
P
Peter Zijlstra 已提交
411
	PN(se.avg_wakeup);
I
Ingo Molnar 已提交
412

413 414
	nr_switches = p->nvcsw + p->nivcsw;

I
Ingo Molnar 已提交
415
#ifdef CONFIG_SCHEDSTATS
I
Ingo Molnar 已提交
416 417 418 419 420 421 422 423
	PN(se.wait_start);
	PN(se.sleep_start);
	PN(se.block_start);
	PN(se.sleep_max);
	PN(se.block_max);
	PN(se.exec_max);
	PN(se.slice_max);
	PN(se.wait_max);
424 425
	PN(se.wait_sum);
	P(se.wait_count);
426 427
	PN(se.iowait_sum);
	P(se.iowait_count);
428
	P(sched_info.bkl_count);
429 430 431 432 433 434 435 436 437 438 439 440 441 442 443 444 445 446 447 448 449 450 451 452 453 454
	P(se.nr_migrations);
	P(se.nr_migrations_cold);
	P(se.nr_failed_migrations_affine);
	P(se.nr_failed_migrations_running);
	P(se.nr_failed_migrations_hot);
	P(se.nr_forced_migrations);
	P(se.nr_wakeups);
	P(se.nr_wakeups_sync);
	P(se.nr_wakeups_migrate);
	P(se.nr_wakeups_local);
	P(se.nr_wakeups_remote);
	P(se.nr_wakeups_affine);
	P(se.nr_wakeups_affine_attempts);
	P(se.nr_wakeups_passive);
	P(se.nr_wakeups_idle);

	{
		u64 avg_atom, avg_per_cpu;

		avg_atom = p->se.sum_exec_runtime;
		if (nr_switches)
			do_div(avg_atom, nr_switches);
		else
			avg_atom = -1LL;

		avg_per_cpu = p->se.sum_exec_runtime;
455
		if (p->se.nr_migrations) {
R
Roman Zippel 已提交
456 457
			avg_per_cpu = div64_u64(avg_per_cpu,
						p->se.nr_migrations);
458
		} else {
459
			avg_per_cpu = -1LL;
460
		}
461 462 463 464

		__PN(avg_atom);
		__PN(avg_per_cpu);
	}
I
Ingo Molnar 已提交
465
#endif
466
	__P(nr_switches);
467
	SEQ_printf(m, "%-35s:%21Ld\n",
468 469 470 471
		   "nr_voluntary_switches", (long long)p->nvcsw);
	SEQ_printf(m, "%-35s:%21Ld\n",
		   "nr_involuntary_switches", (long long)p->nivcsw);

I
Ingo Molnar 已提交
472 473 474
	P(se.load.weight);
	P(policy);
	P(prio);
I
Ingo Molnar 已提交
475
#undef PN
476 477 478
#undef __PN
#undef P
#undef __P
I
Ingo Molnar 已提交
479 480

	{
481
		unsigned int this_cpu = raw_smp_processor_id();
I
Ingo Molnar 已提交
482 483
		u64 t0, t1;

484 485
		t0 = cpu_clock(this_cpu);
		t1 = cpu_clock(this_cpu);
486
		SEQ_printf(m, "%-35s:%21Ld\n",
I
Ingo Molnar 已提交
487 488 489 490 491 492
			   "clock-delta", (long long)(t1-t0));
	}
}

void proc_sched_set_task(struct task_struct *p)
{
I
Ingo Molnar 已提交
493
#ifdef CONFIG_SCHEDSTATS
494
	p->se.wait_max				= 0;
495 496
	p->se.wait_sum				= 0;
	p->se.wait_count			= 0;
497 498
	p->se.iowait_sum			= 0;
	p->se.iowait_count			= 0;
499 500 501 502 503 504 505 506 507 508 509 510 511 512 513 514 515 516 517 518 519
	p->se.sleep_max				= 0;
	p->se.sum_sleep_runtime			= 0;
	p->se.block_max				= 0;
	p->se.exec_max				= 0;
	p->se.slice_max				= 0;
	p->se.nr_migrations			= 0;
	p->se.nr_migrations_cold		= 0;
	p->se.nr_failed_migrations_affine	= 0;
	p->se.nr_failed_migrations_running	= 0;
	p->se.nr_failed_migrations_hot		= 0;
	p->se.nr_forced_migrations		= 0;
	p->se.nr_wakeups			= 0;
	p->se.nr_wakeups_sync			= 0;
	p->se.nr_wakeups_migrate		= 0;
	p->se.nr_wakeups_local			= 0;
	p->se.nr_wakeups_remote			= 0;
	p->se.nr_wakeups_affine			= 0;
	p->se.nr_wakeups_affine_attempts	= 0;
	p->se.nr_wakeups_passive		= 0;
	p->se.nr_wakeups_idle			= 0;
	p->sched_info.bkl_count			= 0;
I
Ingo Molnar 已提交
520
#endif
521 522 523 524
	p->se.sum_exec_runtime			= 0;
	p->se.prev_sum_exec_runtime		= 0;
	p->nvcsw				= 0;
	p->nivcsw				= 0;
I
Ingo Molnar 已提交
525
}