evlist.c 29.1 KB
Newer Older
1 2 3 4 5 6 7 8
/*
 * Copyright (C) 2011, Red Hat Inc, Arnaldo Carvalho de Melo <acme@redhat.com>
 *
 * Parts came from builtin-{top,stat,record}.c, see those files for further
 * copyright notes.
 *
 * Released under the GPL v2. (and only v2, not any later version)
 */
9
#include "util.h"
10
#include <api/fs/debugfs.h>
11
#include <poll.h>
12 13
#include "cpumap.h"
#include "thread_map.h"
14
#include "target.h"
15 16
#include "evlist.h"
#include "evsel.h"
A
Adrian Hunter 已提交
17
#include "debug.h"
18
#include <unistd.h>
19

20
#include "parse-events.h"
21
#include "parse-options.h"
22

23 24
#include <sys/mman.h>

25 26 27
#include <linux/bitops.h>
#include <linux/hash.h>

28
#define FD(e, x, y) (*(int *)xyarray__entry(e->fd, x, y))
29
#define SID(e, x, y) xyarray__entry(e->sample_id, x, y)
30

31 32
void perf_evlist__init(struct perf_evlist *evlist, struct cpu_map *cpus,
		       struct thread_map *threads)
33 34 35 36 37 38
{
	int i;

	for (i = 0; i < PERF_EVLIST__HLIST_SIZE; ++i)
		INIT_HLIST_HEAD(&evlist->heads[i]);
	INIT_LIST_HEAD(&evlist->entries);
39
	perf_evlist__set_maps(evlist, cpus, threads);
40
	evlist->workload.pid = -1;
41 42
}

43
struct perf_evlist *perf_evlist__new(void)
44 45 46
{
	struct perf_evlist *evlist = zalloc(sizeof(*evlist));

47
	if (evlist != NULL)
48
		perf_evlist__init(evlist, NULL, NULL);
49 50 51 52

	return evlist;
}

53 54 55 56 57 58 59 60 61 62 63 64
struct perf_evlist *perf_evlist__new_default(void)
{
	struct perf_evlist *evlist = perf_evlist__new();

	if (evlist && perf_evlist__add_default(evlist)) {
		perf_evlist__delete(evlist);
		evlist = NULL;
	}

	return evlist;
}

65 66 67 68 69 70 71 72 73 74 75 76 77 78 79
/**
 * perf_evlist__set_id_pos - set the positions of event ids.
 * @evlist: selected event list
 *
 * Events with compatible sample types all have the same id_pos
 * and is_pos.  For convenience, put a copy on evlist.
 */
void perf_evlist__set_id_pos(struct perf_evlist *evlist)
{
	struct perf_evsel *first = perf_evlist__first(evlist);

	evlist->id_pos = first->id_pos;
	evlist->is_pos = first->is_pos;
}

80 81 82 83
static void perf_evlist__update_id_pos(struct perf_evlist *evlist)
{
	struct perf_evsel *evsel;

84
	evlist__for_each(evlist, evsel)
85 86 87 88 89
		perf_evsel__calc_id_pos(evsel);

	perf_evlist__set_id_pos(evlist);
}

90 91 92 93
static void perf_evlist__purge(struct perf_evlist *evlist)
{
	struct perf_evsel *pos, *n;

94
	evlist__for_each_safe(evlist, n, pos) {
95 96 97 98 99 100 101
		list_del_init(&pos->node);
		perf_evsel__delete(pos);
	}

	evlist->nr_entries = 0;
}

102
void perf_evlist__exit(struct perf_evlist *evlist)
103
{
104 105
	zfree(&evlist->mmap);
	zfree(&evlist->pollfd);
106 107 108 109
}

void perf_evlist__delete(struct perf_evlist *evlist)
{
110
	perf_evlist__munmap(evlist);
111
	perf_evlist__close(evlist);
112 113 114 115
	cpu_map__delete(evlist->cpus);
	thread_map__delete(evlist->threads);
	evlist->cpus = NULL;
	evlist->threads = NULL;
116 117
	perf_evlist__purge(evlist);
	perf_evlist__exit(evlist);
118 119 120 121 122 123
	free(evlist);
}

void perf_evlist__add(struct perf_evlist *evlist, struct perf_evsel *entry)
{
	list_add_tail(&entry->node, &evlist->entries);
124 125
	entry->idx = evlist->nr_entries;

126 127
	if (!evlist->nr_entries++)
		perf_evlist__set_id_pos(evlist);
128 129
}

130 131 132
void perf_evlist__splice_list_tail(struct perf_evlist *evlist,
				   struct list_head *list,
				   int nr_entries)
133
{
134 135
	bool set_id_pos = !evlist->nr_entries;

136 137
	list_splice_tail(list, &evlist->entries);
	evlist->nr_entries += nr_entries;
138 139
	if (set_id_pos)
		perf_evlist__set_id_pos(evlist);
140 141
}

142 143 144 145 146
void __perf_evlist__set_leader(struct list_head *list)
{
	struct perf_evsel *evsel, *leader;

	leader = list_entry(list->next, struct perf_evsel, node);
147 148 149
	evsel = list_entry(list->prev, struct perf_evsel, node);

	leader->nr_members = evsel->idx - leader->idx + 1;
150

151
	__evlist__for_each(list, evsel) {
152
		evsel->leader = leader;
153 154 155 156
	}
}

void perf_evlist__set_leader(struct perf_evlist *evlist)
157
{
158 159
	if (evlist->nr_entries) {
		evlist->nr_groups = evlist->nr_entries > 1 ? 1 : 0;
160
		__perf_evlist__set_leader(&evlist->entries);
161
	}
162 163
}

164 165 166 167 168 169
int perf_evlist__add_default(struct perf_evlist *evlist)
{
	struct perf_event_attr attr = {
		.type = PERF_TYPE_HARDWARE,
		.config = PERF_COUNT_HW_CPU_CYCLES,
	};
170 171 172
	struct perf_evsel *evsel;

	event_attr_init(&attr);
173

174
	evsel = perf_evsel__new(&attr);
175
	if (evsel == NULL)
176 177 178 179 180 181
		goto error;

	/* use strdup() because free(evsel) assumes name is allocated */
	evsel->name = strdup("cycles");
	if (!evsel->name)
		goto error_free;
182 183 184

	perf_evlist__add(evlist, evsel);
	return 0;
185 186 187 188
error_free:
	perf_evsel__delete(evsel);
error:
	return -ENOMEM;
189
}
190

191 192
static int perf_evlist__add_attrs(struct perf_evlist *evlist,
				  struct perf_event_attr *attrs, size_t nr_attrs)
193 194 195 196 197 198
{
	struct perf_evsel *evsel, *n;
	LIST_HEAD(head);
	size_t i;

	for (i = 0; i < nr_attrs; i++) {
199
		evsel = perf_evsel__new_idx(attrs + i, evlist->nr_entries + i);
200 201 202 203 204 205 206 207 208 209
		if (evsel == NULL)
			goto out_delete_partial_list;
		list_add_tail(&evsel->node, &head);
	}

	perf_evlist__splice_list_tail(evlist, &head, nr_attrs);

	return 0;

out_delete_partial_list:
210
	__evlist__for_each_safe(&head, n, evsel)
211 212 213 214
		perf_evsel__delete(evsel);
	return -1;
}

215 216 217 218 219 220 221 222 223 224 225
int __perf_evlist__add_default_attrs(struct perf_evlist *evlist,
				     struct perf_event_attr *attrs, size_t nr_attrs)
{
	size_t i;

	for (i = 0; i < nr_attrs; i++)
		event_attr_init(attrs + i);

	return perf_evlist__add_attrs(evlist, attrs, nr_attrs);
}

226 227
struct perf_evsel *
perf_evlist__find_tracepoint_by_id(struct perf_evlist *evlist, int id)
228 229 230
{
	struct perf_evsel *evsel;

231
	evlist__for_each(evlist, evsel) {
232 233 234 235 236 237 238 239
		if (evsel->attr.type   == PERF_TYPE_TRACEPOINT &&
		    (int)evsel->attr.config == id)
			return evsel;
	}

	return NULL;
}

240 241 242 243 244 245
struct perf_evsel *
perf_evlist__find_tracepoint_by_name(struct perf_evlist *evlist,
				     const char *name)
{
	struct perf_evsel *evsel;

246
	evlist__for_each(evlist, evsel) {
247 248 249 250 251 252 253 254
		if ((evsel->attr.type == PERF_TYPE_TRACEPOINT) &&
		    (strcmp(evsel->name, name) == 0))
			return evsel;
	}

	return NULL;
}

255 256 257
int perf_evlist__add_newtp(struct perf_evlist *evlist,
			   const char *sys, const char *name, void *handler)
{
258
	struct perf_evsel *evsel = perf_evsel__newtp(sys, name);
259 260 261 262

	if (evsel == NULL)
		return -1;

263
	evsel->handler = handler;
264 265 266 267
	perf_evlist__add(evlist, evsel);
	return 0;
}

268 269 270 271 272 273 274 275 276
static int perf_evlist__nr_threads(struct perf_evlist *evlist,
				   struct perf_evsel *evsel)
{
	if (evsel->system_wide)
		return 1;
	else
		return thread_map__nr(evlist->threads);
}

277 278 279 280
void perf_evlist__disable(struct perf_evlist *evlist)
{
	int cpu, thread;
	struct perf_evsel *pos;
281
	int nr_cpus = cpu_map__nr(evlist->cpus);
282
	int nr_threads;
283

284
	for (cpu = 0; cpu < nr_cpus; cpu++) {
285
		evlist__for_each(evlist, pos) {
286
			if (!perf_evsel__is_group_leader(pos) || !pos->fd)
287
				continue;
288
			nr_threads = perf_evlist__nr_threads(evlist, pos);
289
			for (thread = 0; thread < nr_threads; thread++)
290 291
				ioctl(FD(pos, cpu, thread),
				      PERF_EVENT_IOC_DISABLE, 0);
292 293 294 295
		}
	}
}

296 297 298 299
void perf_evlist__enable(struct perf_evlist *evlist)
{
	int cpu, thread;
	struct perf_evsel *pos;
300
	int nr_cpus = cpu_map__nr(evlist->cpus);
301
	int nr_threads;
302

303
	for (cpu = 0; cpu < nr_cpus; cpu++) {
304
		evlist__for_each(evlist, pos) {
305
			if (!perf_evsel__is_group_leader(pos) || !pos->fd)
306
				continue;
307
			nr_threads = perf_evlist__nr_threads(evlist, pos);
308
			for (thread = 0; thread < nr_threads; thread++)
309 310
				ioctl(FD(pos, cpu, thread),
				      PERF_EVENT_IOC_ENABLE, 0);
311 312 313 314
		}
	}
}

315 316 317 318
int perf_evlist__disable_event(struct perf_evlist *evlist,
			       struct perf_evsel *evsel)
{
	int cpu, thread, err;
319 320
	int nr_cpus = cpu_map__nr(evlist->cpus);
	int nr_threads = perf_evlist__nr_threads(evlist, evsel);
321 322 323 324

	if (!evsel->fd)
		return 0;

325 326
	for (cpu = 0; cpu < nr_cpus; cpu++) {
		for (thread = 0; thread < nr_threads; thread++) {
327 328 329 330 331 332 333 334 335 336 337 338 339
			err = ioctl(FD(evsel, cpu, thread),
				    PERF_EVENT_IOC_DISABLE, 0);
			if (err)
				return err;
		}
	}
	return 0;
}

int perf_evlist__enable_event(struct perf_evlist *evlist,
			      struct perf_evsel *evsel)
{
	int cpu, thread, err;
340 341
	int nr_cpus = cpu_map__nr(evlist->cpus);
	int nr_threads = perf_evlist__nr_threads(evlist, evsel);
342 343 344 345

	if (!evsel->fd)
		return -EINVAL;

346 347
	for (cpu = 0; cpu < nr_cpus; cpu++) {
		for (thread = 0; thread < nr_threads; thread++) {
348 349 350 351 352 353 354 355 356
			err = ioctl(FD(evsel, cpu, thread),
				    PERF_EVENT_IOC_ENABLE, 0);
			if (err)
				return err;
		}
	}
	return 0;
}

357
static int perf_evlist__alloc_pollfd(struct perf_evlist *evlist)
358
{
359 360
	int nr_cpus = cpu_map__nr(evlist->cpus);
	int nr_threads = thread_map__nr(evlist->threads);
361 362 363 364 365 366 367 368 369 370
	int nfds = 0;
	struct perf_evsel *evsel;

	list_for_each_entry(evsel, &evlist->entries, node) {
		if (evsel->system_wide)
			nfds += nr_cpus;
		else
			nfds += nr_cpus * nr_threads;
	}

371 372 373
	evlist->pollfd = malloc(sizeof(struct pollfd) * nfds);
	return evlist->pollfd != NULL ? 0 : -ENOMEM;
}
374 375 376 377 378 379 380 381

void perf_evlist__add_pollfd(struct perf_evlist *evlist, int fd)
{
	fcntl(fd, F_SETFL, O_NONBLOCK);
	evlist->pollfd[evlist->nr_fds].fd = fd;
	evlist->pollfd[evlist->nr_fds].events = POLLIN;
	evlist->nr_fds++;
}
382

383 384 385
static void perf_evlist__id_hash(struct perf_evlist *evlist,
				 struct perf_evsel *evsel,
				 int cpu, int thread, u64 id)
386 387 388 389 390 391 392 393 394 395
{
	int hash;
	struct perf_sample_id *sid = SID(evsel, cpu, thread);

	sid->id = id;
	sid->evsel = evsel;
	hash = hash_64(sid->id, PERF_EVLIST__HLIST_BITS);
	hlist_add_head(&sid->node, &evlist->heads[hash]);
}

396 397 398 399 400 401 402 403 404 405
void perf_evlist__id_add(struct perf_evlist *evlist, struct perf_evsel *evsel,
			 int cpu, int thread, u64 id)
{
	perf_evlist__id_hash(evlist, evsel, cpu, thread, id);
	evsel->id[evsel->ids++] = id;
}

static int perf_evlist__id_add_fd(struct perf_evlist *evlist,
				  struct perf_evsel *evsel,
				  int cpu, int thread, int fd)
406 407
{
	u64 read_data[4] = { 0, };
408
	int id_idx = 1; /* The first entry is the counter value */
409 410 411 412 413 414 415 416 417 418 419
	u64 id;
	int ret;

	ret = ioctl(fd, PERF_EVENT_IOC_ID, &id);
	if (!ret)
		goto add;

	if (errno != ENOTTY)
		return -1;

	/* Legacy way to get event id.. All hail to old kernels! */
420

421 422 423 424 425 426 427
	/*
	 * This way does not work with group format read, so bail
	 * out in that case.
	 */
	if (perf_evlist__read_format(evlist) & PERF_FORMAT_GROUP)
		return -1;

428 429 430 431 432 433 434 435 436
	if (!(evsel->attr.read_format & PERF_FORMAT_ID) ||
	    read(fd, &read_data, sizeof(read_data)) == -1)
		return -1;

	if (evsel->attr.read_format & PERF_FORMAT_TOTAL_TIME_ENABLED)
		++id_idx;
	if (evsel->attr.read_format & PERF_FORMAT_TOTAL_TIME_RUNNING)
		++id_idx;

437 438 439 440
	id = read_data[id_idx];

 add:
	perf_evlist__id_add(evlist, evsel, cpu, thread, id);
441 442 443
	return 0;
}

444
struct perf_sample_id *perf_evlist__id2sid(struct perf_evlist *evlist, u64 id)
445 446 447 448 449 450 451 452
{
	struct hlist_head *head;
	struct perf_sample_id *sid;
	int hash;

	hash = hash_64(id, PERF_EVLIST__HLIST_BITS);
	head = &evlist->heads[hash];

453
	hlist_for_each_entry(sid, head, node)
454
		if (sid->id == id)
455 456 457 458 459 460 461 462 463 464 465 466 467 468 469
			return sid;

	return NULL;
}

struct perf_evsel *perf_evlist__id2evsel(struct perf_evlist *evlist, u64 id)
{
	struct perf_sample_id *sid;

	if (evlist->nr_entries == 1)
		return perf_evlist__first(evlist);

	sid = perf_evlist__id2sid(evlist, id);
	if (sid)
		return sid->evsel;
470 471

	if (!perf_evlist__sample_id_all(evlist))
472
		return perf_evlist__first(evlist);
473

474 475
	return NULL;
}
476

477 478 479 480 481 482 483 484 485 486 487 488 489 490 491 492 493 494 495 496 497 498 499 500
static int perf_evlist__event2id(struct perf_evlist *evlist,
				 union perf_event *event, u64 *id)
{
	const u64 *array = event->sample.array;
	ssize_t n;

	n = (event->header.size - sizeof(event->header)) >> 3;

	if (event->header.type == PERF_RECORD_SAMPLE) {
		if (evlist->id_pos >= n)
			return -1;
		*id = array[evlist->id_pos];
	} else {
		if (evlist->is_pos > n)
			return -1;
		n -= evlist->is_pos;
		*id = array[n];
	}
	return 0;
}

static struct perf_evsel *perf_evlist__event2evsel(struct perf_evlist *evlist,
						   union perf_event *event)
{
501
	struct perf_evsel *first = perf_evlist__first(evlist);
502 503 504 505 506 507
	struct hlist_head *head;
	struct perf_sample_id *sid;
	int hash;
	u64 id;

	if (evlist->nr_entries == 1)
508 509 510 511 512
		return first;

	if (!first->attr.sample_id_all &&
	    event->header.type != PERF_RECORD_SAMPLE)
		return first;
513 514 515 516 517 518

	if (perf_evlist__event2id(evlist, event, &id))
		return NULL;

	/* Synthesized events have an id of zero */
	if (!id)
519
		return first;
520 521 522 523 524 525 526 527 528 529 530

	hash = hash_64(id, PERF_EVLIST__HLIST_BITS);
	head = &evlist->heads[hash];

	hlist_for_each_entry(sid, head, node) {
		if (sid->id == id)
			return sid->evsel;
	}
	return NULL;
}

531
union perf_event *perf_evlist__mmap_read(struct perf_evlist *evlist, int idx)
532
{
533
	struct perf_mmap *md = &evlist->mmap[idx];
534 535 536
	unsigned int head = perf_mmap__read_head(md);
	unsigned int old = md->prev;
	unsigned char *data = md->base + page_size;
537
	union perf_event *event = NULL;
538

539
	if (evlist->overwrite) {
540
		/*
541 542 543 544 545 546
		 * If we're further behind than half the buffer, there's a chance
		 * the writer will bite our tail and mess up the samples under us.
		 *
		 * If we somehow ended up ahead of the head, we got messed up.
		 *
		 * In either case, truncate and restart at head.
547
		 */
548 549 550 551 552 553 554 555 556
		int diff = head - old;
		if (diff > md->mask / 2 || diff < 0) {
			fprintf(stderr, "WARNING: failed to keep up with mmap data.\n");

			/*
			 * head points to a known good entry, start there.
			 */
			old = head;
		}
557 558 559 560 561
	}

	if (old != head) {
		size_t size;

562
		event = (union perf_event *)&data[old & md->mask];
563 564 565 566 567 568 569 570 571
		size = event->header.size;

		/*
		 * Event straddles the mmap boundary -- header should always
		 * be inside due to u64 alignment of output.
		 */
		if ((old & md->mask) + size != ((old + size) & md->mask)) {
			unsigned int offset = old;
			unsigned int len = min(sizeof(*event), size), cpy;
572
			void *dst = md->event_copy;
573 574 575 576 577 578 579 580 581

			do {
				cpy = min(md->mask + 1 - (offset & md->mask), len);
				memcpy(dst, &data[offset & md->mask], cpy);
				offset += cpy;
				dst += cpy;
				len -= cpy;
			} while (len);

582
			event = (union perf_event *) md->event_copy;
583 584 585 586 587 588
		}

		old += size;
	}

	md->prev = old;
589

590 591
	return event;
}
592

593 594 595 596 597 598 599 600 601 602
void perf_evlist__mmap_consume(struct perf_evlist *evlist, int idx)
{
	if (!evlist->overwrite) {
		struct perf_mmap *md = &evlist->mmap[idx];
		unsigned int old = md->prev;

		perf_mmap__write_tail(md, old);
	}
}

603 604 605 606 607 608 609 610
static void __perf_evlist__munmap(struct perf_evlist *evlist, int idx)
{
	if (evlist->mmap[idx].base != NULL) {
		munmap(evlist->mmap[idx].base, evlist->mmap_len);
		evlist->mmap[idx].base = NULL;
	}
}

611
void perf_evlist__munmap(struct perf_evlist *evlist)
612
{
613
	int i;
614

615 616 617
	if (evlist->mmap == NULL)
		return;

618 619
	for (i = 0; i < evlist->nr_mmaps; i++)
		__perf_evlist__munmap(evlist, i);
620

621
	zfree(&evlist->mmap);
622 623
}

624
static int perf_evlist__alloc_mmap(struct perf_evlist *evlist)
625
{
626
	evlist->nr_mmaps = cpu_map__nr(evlist->cpus);
627
	if (cpu_map__empty(evlist->cpus))
628
		evlist->nr_mmaps = thread_map__nr(evlist->threads);
629
	evlist->mmap = zalloc(evlist->nr_mmaps * sizeof(struct perf_mmap));
630 631 632
	return evlist->mmap != NULL ? 0 : -ENOMEM;
}

633 634 635 636 637 638 639
struct mmap_params {
	int prot;
	int mask;
};

static int __perf_evlist__mmap(struct perf_evlist *evlist, int idx,
			       struct mmap_params *mp, int fd)
640
{
641
	evlist->mmap[idx].prev = 0;
642 643
	evlist->mmap[idx].mask = mp->mask;
	evlist->mmap[idx].base = mmap(NULL, evlist->mmap_len, mp->prot,
644
				      MAP_SHARED, fd, 0);
645
	if (evlist->mmap[idx].base == MAP_FAILED) {
646 647
		pr_debug2("failed to mmap perf event ring buffer, error %d\n",
			  errno);
648
		evlist->mmap[idx].base = NULL;
649
		return -1;
650
	}
651 652 653 654 655

	perf_evlist__add_pollfd(evlist, fd);
	return 0;
}

656
static int perf_evlist__mmap_per_evsel(struct perf_evlist *evlist, int idx,
657 658
				       struct mmap_params *mp, int cpu,
				       int thread, int *output)
659 660
{
	struct perf_evsel *evsel;
661

662
	evlist__for_each(evlist, evsel) {
663 664 665 666 667 668
		int fd;

		if (evsel->system_wide && thread)
			continue;

		fd = FD(evsel, cpu, thread);
669 670 671

		if (*output == -1) {
			*output = fd;
672
			if (__perf_evlist__mmap(evlist, idx, mp, *output) < 0)
673 674 675 676 677 678 679 680 681 682 683 684 685 686
				return -1;
		} else {
			if (ioctl(fd, PERF_EVENT_IOC_SET_OUTPUT, *output) != 0)
				return -1;
		}

		if ((evsel->attr.read_format & PERF_FORMAT_ID) &&
		    perf_evlist__id_add_fd(evlist, evsel, cpu, thread, fd) < 0)
			return -1;
	}

	return 0;
}

687 688
static int perf_evlist__mmap_per_cpu(struct perf_evlist *evlist,
				     struct mmap_params *mp)
689
{
690
	int cpu, thread;
691 692
	int nr_cpus = cpu_map__nr(evlist->cpus);
	int nr_threads = thread_map__nr(evlist->threads);
693

A
Adrian Hunter 已提交
694
	pr_debug2("perf event ring buffer mmapped per cpu\n");
695
	for (cpu = 0; cpu < nr_cpus; cpu++) {
696 697
		int output = -1;

698
		for (thread = 0; thread < nr_threads; thread++) {
699 700
			if (perf_evlist__mmap_per_evsel(evlist, cpu, mp, cpu,
							thread, &output))
701
				goto out_unmap;
702 703 704 705 706 707
		}
	}

	return 0;

out_unmap:
708 709
	for (cpu = 0; cpu < nr_cpus; cpu++)
		__perf_evlist__munmap(evlist, cpu);
710 711 712
	return -1;
}

713 714
static int perf_evlist__mmap_per_thread(struct perf_evlist *evlist,
					struct mmap_params *mp)
715 716
{
	int thread;
717
	int nr_threads = thread_map__nr(evlist->threads);
718

A
Adrian Hunter 已提交
719
	pr_debug2("perf event ring buffer mmapped per thread\n");
720
	for (thread = 0; thread < nr_threads; thread++) {
721 722
		int output = -1;

723 724
		if (perf_evlist__mmap_per_evsel(evlist, thread, mp, 0, thread,
						&output))
725
			goto out_unmap;
726 727 728 729 730
	}

	return 0;

out_unmap:
731 732
	for (thread = 0; thread < nr_threads; thread++)
		__perf_evlist__munmap(evlist, thread);
733 734 735
	return -1;
}

736 737 738 739 740 741 742 743 744 745 746
static size_t perf_evlist__mmap_size(unsigned long pages)
{
	/* 512 kiB: default amount of unprivileged mlocked memory */
	if (pages == UINT_MAX)
		pages = (512 * 1024) / page_size;
	else if (!is_power_of_2(pages))
		return 0;

	return (pages + 1) * page_size;
}

747 748
static long parse_pages_arg(const char *str, unsigned long min,
			    unsigned long max)
749
{
750
	unsigned long pages, val;
751 752 753 754 755 756 757
	static struct parse_tag tags[] = {
		{ .tag  = 'B', .mult = 1       },
		{ .tag  = 'K', .mult = 1 << 10 },
		{ .tag  = 'M', .mult = 1 << 20 },
		{ .tag  = 'G', .mult = 1 << 30 },
		{ .tag  = 0 },
	};
758

759
	if (str == NULL)
760
		return -EINVAL;
761

762
	val = parse_tag_value(str, tags);
763
	if (val != (unsigned long) -1) {
764 765 766 767 768 769
		/* we got file size value */
		pages = PERF_ALIGN(val, page_size) / page_size;
	} else {
		/* we got pages count value */
		char *eptr;
		pages = strtoul(str, &eptr, 10);
770 771
		if (*eptr != '\0')
			return -EINVAL;
772 773
	}

774
	if (pages == 0 && min == 0) {
775
		/* leave number of pages at 0 */
776
	} else if (!is_power_of_2(pages)) {
777
		/* round pages up to next power of 2 */
778 779 780
		pages = next_pow2_l(pages);
		if (!pages)
			return -EINVAL;
781 782
		pr_info("rounding mmap pages size to %lu bytes (%lu pages)\n",
			pages * page_size, pages);
783 784
	}

785 786 787 788 789 790 791 792 793 794 795 796 797
	if (pages > max)
		return -EINVAL;

	return pages;
}

int perf_evlist__parse_mmap_pages(const struct option *opt, const char *str,
				  int unset __maybe_unused)
{
	unsigned int *mmap_pages = opt->value;
	unsigned long max = UINT_MAX;
	long pages;

A
Adrian Hunter 已提交
798
	if (max > SIZE_MAX / page_size)
799 800 801 802 803
		max = SIZE_MAX / page_size;

	pages = parse_pages_arg(str, 1, max);
	if (pages < 0) {
		pr_err("Invalid argument for --mmap_pages/-m\n");
804 805 806 807 808 809 810
		return -1;
	}

	*mmap_pages = pages;
	return 0;
}

811 812 813 814 815
/**
 * perf_evlist__mmap - Create mmaps to receive events.
 * @evlist: list of events
 * @pages: map length in pages
 * @overwrite: overwrite older events?
816
 *
817 818 819
 * If @overwrite is %false the user needs to signal event consumption using
 * perf_mmap__write_tail().  Using perf_evlist__mmap_read() does this
 * automatically.
820
 *
821
 * Return: %0 on success, negative error code otherwise.
822
 */
823 824
int perf_evlist__mmap(struct perf_evlist *evlist, unsigned int pages,
		      bool overwrite)
825
{
826
	struct perf_evsel *evsel;
827 828
	const struct cpu_map *cpus = evlist->cpus;
	const struct thread_map *threads = evlist->threads;
829 830 831
	struct mmap_params mp = {
		.prot = PROT_READ | (overwrite ? 0 : PROT_WRITE),
	};
832

833
	if (evlist->mmap == NULL && perf_evlist__alloc_mmap(evlist) < 0)
834 835
		return -ENOMEM;

836
	if (evlist->pollfd == NULL && perf_evlist__alloc_pollfd(evlist) < 0)
837 838 839
		return -ENOMEM;

	evlist->overwrite = overwrite;
840
	evlist->mmap_len = perf_evlist__mmap_size(pages);
841
	pr_debug("mmap size %zuB\n", evlist->mmap_len);
842
	mp.mask = evlist->mmap_len - page_size - 1;
843

844
	evlist__for_each(evlist, evsel) {
845
		if ((evsel->attr.read_format & PERF_FORMAT_ID) &&
846
		    evsel->sample_id == NULL &&
847
		    perf_evsel__alloc_id(evsel, cpu_map__nr(cpus), threads->nr) < 0)
848 849 850
			return -ENOMEM;
	}

851
	if (cpu_map__empty(cpus))
852
		return perf_evlist__mmap_per_thread(evlist, &mp);
853

854
	return perf_evlist__mmap_per_cpu(evlist, &mp);
855
}
856

857
int perf_evlist__create_maps(struct perf_evlist *evlist, struct target *target)
858
{
859 860
	evlist->threads = thread_map__new_str(target->pid, target->tid,
					      target->uid);
861 862 863 864

	if (evlist->threads == NULL)
		return -1;

865
	if (target__uses_dummy_map(target))
N
Namhyung Kim 已提交
866
		evlist->cpus = cpu_map__dummy_new();
867 868
	else
		evlist->cpus = cpu_map__new(target->cpu_list);
869 870 871 872 873 874 875 876 877 878 879

	if (evlist->cpus == NULL)
		goto out_delete_threads;

	return 0;

out_delete_threads:
	thread_map__delete(evlist->threads);
	return -1;
}

880
int perf_evlist__apply_filters(struct perf_evlist *evlist)
881 882
{
	struct perf_evsel *evsel;
883 884
	int err = 0;
	const int ncpus = cpu_map__nr(evlist->cpus),
885
		  nthreads = thread_map__nr(evlist->threads);
886

887
	evlist__for_each(evlist, evsel) {
888
		if (evsel->filter == NULL)
889
			continue;
890 891 892 893

		err = perf_evsel__set_filter(evsel, ncpus, nthreads, evsel->filter);
		if (err)
			break;
894 895
	}

896 897 898 899 900 901 902 903
	return err;
}

int perf_evlist__set_filter(struct perf_evlist *evlist, const char *filter)
{
	struct perf_evsel *evsel;
	int err = 0;
	const int ncpus = cpu_map__nr(evlist->cpus),
904
		  nthreads = thread_map__nr(evlist->threads);
905

906
	evlist__for_each(evlist, evsel) {
907 908 909 910 911 912
		err = perf_evsel__set_filter(evsel, ncpus, nthreads, filter);
		if (err)
			break;
	}

	return err;
913
}
914

915
bool perf_evlist__valid_sample_type(struct perf_evlist *evlist)
916
{
917
	struct perf_evsel *pos;
918

919 920 921 922 923 924
	if (evlist->nr_entries == 1)
		return true;

	if (evlist->id_pos < 0 || evlist->is_pos < 0)
		return false;

925
	evlist__for_each(evlist, pos) {
926 927
		if (pos->id_pos != evlist->id_pos ||
		    pos->is_pos != evlist->is_pos)
928
			return false;
929 930
	}

931
	return true;
932 933
}

934
u64 __perf_evlist__combined_sample_type(struct perf_evlist *evlist)
935
{
936 937 938 939 940
	struct perf_evsel *evsel;

	if (evlist->combined_sample_type)
		return evlist->combined_sample_type;

941
	evlist__for_each(evlist, evsel)
942 943 944 945 946 947 948 949 950
		evlist->combined_sample_type |= evsel->attr.sample_type;

	return evlist->combined_sample_type;
}

u64 perf_evlist__combined_sample_type(struct perf_evlist *evlist)
{
	evlist->combined_sample_type = 0;
	return __perf_evlist__combined_sample_type(evlist);
951 952
}

953 954 955 956 957 958
bool perf_evlist__valid_read_format(struct perf_evlist *evlist)
{
	struct perf_evsel *first = perf_evlist__first(evlist), *pos = first;
	u64 read_format = first->attr.read_format;
	u64 sample_type = first->attr.sample_type;

959
	evlist__for_each(evlist, pos) {
960 961 962 963 964 965 966 967 968 969 970 971 972 973 974 975 976 977 978
		if (read_format != pos->attr.read_format)
			return false;
	}

	/* PERF_SAMPLE_READ imples PERF_FORMAT_ID. */
	if ((sample_type & PERF_SAMPLE_READ) &&
	    !(read_format & PERF_FORMAT_ID)) {
		return false;
	}

	return true;
}

u64 perf_evlist__read_format(struct perf_evlist *evlist)
{
	struct perf_evsel *first = perf_evlist__first(evlist);
	return first->attr.read_format;
}

979
u16 perf_evlist__id_hdr_size(struct perf_evlist *evlist)
980
{
981
	struct perf_evsel *first = perf_evlist__first(evlist);
982 983 984 985 986 987 988 989 990 991 992 993 994 995 996 997 998 999 1000 1001 1002 1003 1004
	struct perf_sample *data;
	u64 sample_type;
	u16 size = 0;

	if (!first->attr.sample_id_all)
		goto out;

	sample_type = first->attr.sample_type;

	if (sample_type & PERF_SAMPLE_TID)
		size += sizeof(data->tid) * 2;

       if (sample_type & PERF_SAMPLE_TIME)
		size += sizeof(data->time);

	if (sample_type & PERF_SAMPLE_ID)
		size += sizeof(data->id);

	if (sample_type & PERF_SAMPLE_STREAM_ID)
		size += sizeof(data->stream_id);

	if (sample_type & PERF_SAMPLE_CPU)
		size += sizeof(data->cpu) * 2;
1005 1006 1007

	if (sample_type & PERF_SAMPLE_IDENTIFIER)
		size += sizeof(data->id);
1008 1009 1010 1011
out:
	return size;
}

1012
bool perf_evlist__valid_sample_id_all(struct perf_evlist *evlist)
1013
{
1014
	struct perf_evsel *first = perf_evlist__first(evlist), *pos = first;
1015

1016
	evlist__for_each_continue(evlist, pos) {
1017 1018
		if (first->attr.sample_id_all != pos->attr.sample_id_all)
			return false;
1019 1020
	}

1021 1022 1023
	return true;
}

1024
bool perf_evlist__sample_id_all(struct perf_evlist *evlist)
1025
{
1026
	struct perf_evsel *first = perf_evlist__first(evlist);
1027
	return first->attr.sample_id_all;
1028
}
1029 1030 1031 1032 1033 1034

void perf_evlist__set_selected(struct perf_evlist *evlist,
			       struct perf_evsel *evsel)
{
	evlist->selected = evsel;
}
1035

1036 1037 1038 1039 1040
void perf_evlist__close(struct perf_evlist *evlist)
{
	struct perf_evsel *evsel;
	int ncpus = cpu_map__nr(evlist->cpus);
	int nthreads = thread_map__nr(evlist->threads);
1041
	int n;
1042

1043 1044 1045 1046
	evlist__for_each_reverse(evlist, evsel) {
		n = evsel->cpus ? evsel->cpus->nr : ncpus;
		perf_evsel__close(evsel, n, nthreads);
	}
1047 1048
}

1049
int perf_evlist__open(struct perf_evlist *evlist)
1050
{
1051
	struct perf_evsel *evsel;
1052
	int err;
1053

1054 1055
	perf_evlist__update_id_pos(evlist);

1056
	evlist__for_each(evlist, evsel) {
1057
		err = perf_evsel__open(evsel, evlist->cpus, evlist->threads);
1058 1059 1060 1061 1062 1063
		if (err < 0)
			goto out_err;
	}

	return 0;
out_err:
1064
	perf_evlist__close(evlist);
1065
	errno = -err;
1066 1067
	return err;
}
1068

1069
int perf_evlist__prepare_workload(struct perf_evlist *evlist, struct target *target,
1070
				  const char *argv[], bool pipe_output,
1071
				  void (*exec_error)(int signo, siginfo_t *info, void *ucontext))
1072 1073 1074 1075 1076 1077 1078 1079 1080 1081 1082 1083 1084 1085 1086 1087 1088 1089 1090 1091 1092
{
	int child_ready_pipe[2], go_pipe[2];
	char bf;

	if (pipe(child_ready_pipe) < 0) {
		perror("failed to create 'ready' pipe");
		return -1;
	}

	if (pipe(go_pipe) < 0) {
		perror("failed to create 'go' pipe");
		goto out_close_ready_pipe;
	}

	evlist->workload.pid = fork();
	if (evlist->workload.pid < 0) {
		perror("failed to fork");
		goto out_close_pipes;
	}

	if (!evlist->workload.pid) {
1093 1094
		int ret;

1095
		if (pipe_output)
1096 1097
			dup2(2, 1);

1098 1099
		signal(SIGTERM, SIG_DFL);

1100 1101 1102 1103 1104 1105 1106 1107 1108 1109 1110 1111
		close(child_ready_pipe[0]);
		close(go_pipe[1]);
		fcntl(go_pipe[0], F_SETFD, FD_CLOEXEC);

		/*
		 * Tell the parent we're ready to go
		 */
		close(child_ready_pipe[1]);

		/*
		 * Wait until the parent tells us to go.
		 */
1112 1113 1114 1115 1116 1117 1118 1119 1120 1121 1122 1123 1124 1125 1126 1127
		ret = read(go_pipe[0], &bf, 1);
		/*
		 * The parent will ask for the execvp() to be performed by
		 * writing exactly one byte, in workload.cork_fd, usually via
		 * perf_evlist__start_workload().
		 *
		 * For cancelling the workload without actuallin running it,
		 * the parent will just close workload.cork_fd, without writing
		 * anything, i.e. read will return zero and we just exit()
		 * here.
		 */
		if (ret != 1) {
			if (ret == -1)
				perror("unable to read pipe");
			exit(ret);
		}
1128 1129 1130

		execvp(argv[0], (char **)argv);

1131
		if (exec_error) {
1132 1133 1134 1135 1136 1137 1138
			union sigval val;

			val.sival_int = errno;
			if (sigqueue(getppid(), SIGUSR1, val))
				perror(argv[0]);
		} else
			perror(argv[0]);
1139 1140 1141
		exit(-1);
	}

1142 1143 1144 1145 1146 1147 1148 1149
	if (exec_error) {
		struct sigaction act = {
			.sa_flags     = SA_SIGINFO,
			.sa_sigaction = exec_error,
		};
		sigaction(SIGUSR1, &act, NULL);
	}

1150
	if (target__none(target))
1151 1152 1153 1154 1155 1156 1157 1158 1159 1160 1161 1162
		evlist->threads->map[0] = evlist->workload.pid;

	close(child_ready_pipe[1]);
	close(go_pipe[0]);
	/*
	 * wait for child to settle
	 */
	if (read(child_ready_pipe[0], &bf, 1) == -1) {
		perror("unable to read pipe");
		goto out_close_pipes;
	}

1163
	fcntl(go_pipe[1], F_SETFD, FD_CLOEXEC);
1164 1165 1166 1167 1168 1169 1170 1171 1172 1173 1174 1175 1176 1177 1178 1179
	evlist->workload.cork_fd = go_pipe[1];
	close(child_ready_pipe[0]);
	return 0;

out_close_pipes:
	close(go_pipe[0]);
	close(go_pipe[1]);
out_close_ready_pipe:
	close(child_ready_pipe[0]);
	close(child_ready_pipe[1]);
	return -1;
}

int perf_evlist__start_workload(struct perf_evlist *evlist)
{
	if (evlist->workload.cork_fd > 0) {
1180
		char bf = 0;
1181
		int ret;
1182 1183 1184
		/*
		 * Remove the cork, let it rip!
		 */
1185 1186 1187 1188 1189 1190
		ret = write(evlist->workload.cork_fd, &bf, 1);
		if (ret < 0)
			perror("enable to write to pipe");

		close(evlist->workload.cork_fd);
		return ret;
1191 1192 1193 1194
	}

	return 0;
}
1195

1196
int perf_evlist__parse_sample(struct perf_evlist *evlist, union perf_event *event,
1197
			      struct perf_sample *sample)
1198
{
1199 1200 1201 1202
	struct perf_evsel *evsel = perf_evlist__event2evsel(evlist, event);

	if (!evsel)
		return -EFAULT;
1203
	return perf_evsel__parse_sample(evsel, event, sample);
1204
}
1205 1206 1207 1208 1209 1210

size_t perf_evlist__fprintf(struct perf_evlist *evlist, FILE *fp)
{
	struct perf_evsel *evsel;
	size_t printed = 0;

1211
	evlist__for_each(evlist, evsel) {
1212 1213 1214 1215
		printed += fprintf(fp, "%s%s", evsel->idx ? ", " : "",
				   perf_evsel__name(evsel));
	}

1216
	return printed + fprintf(fp, "\n");
1217
}
1218 1219 1220 1221 1222 1223 1224 1225 1226 1227 1228 1229 1230 1231 1232 1233 1234 1235 1236 1237 1238 1239 1240 1241 1242 1243 1244

int perf_evlist__strerror_tp(struct perf_evlist *evlist __maybe_unused,
			     int err, char *buf, size_t size)
{
	char sbuf[128];

	switch (err) {
	case ENOENT:
		scnprintf(buf, size, "%s",
			  "Error:\tUnable to find debugfs\n"
			  "Hint:\tWas your kernel was compiled with debugfs support?\n"
			  "Hint:\tIs the debugfs filesystem mounted?\n"
			  "Hint:\tTry 'sudo mount -t debugfs nodev /sys/kernel/debug'");
		break;
	case EACCES:
		scnprintf(buf, size,
			  "Error:\tNo permissions to read %s/tracing/events/raw_syscalls\n"
			  "Hint:\tTry 'sudo mount -o remount,mode=755 %s'\n",
			  debugfs_mountpoint, debugfs_mountpoint);
		break;
	default:
		scnprintf(buf, size, "%s", strerror_r(err, sbuf, sizeof(sbuf)));
		break;
	}

	return 0;
}
1245 1246 1247 1248 1249 1250 1251 1252 1253 1254 1255 1256 1257 1258

int perf_evlist__strerror_open(struct perf_evlist *evlist __maybe_unused,
			       int err, char *buf, size_t size)
{
	int printed, value;
	char sbuf[128], *emsg = strerror_r(err, sbuf, sizeof(sbuf));

	switch (err) {
	case EACCES:
	case EPERM:
		printed = scnprintf(buf, size,
				    "Error:\t%s.\n"
				    "Hint:\tCheck /proc/sys/kernel/perf_event_paranoid setting.", emsg);

1259
		value = perf_event_paranoid();
1260 1261 1262 1263 1264 1265 1266 1267

		printed += scnprintf(buf + printed, size - printed, "\nHint:\t");

		if (value >= 2) {
			printed += scnprintf(buf + printed, size - printed,
					     "For your workloads it needs to be <= 1\nHint:\t");
		}
		printed += scnprintf(buf + printed, size - printed,
1268
				     "For system wide tracing it needs to be set to -1.\n");
1269 1270

		printed += scnprintf(buf + printed, size - printed,
1271 1272
				    "Hint:\tTry: 'sudo sh -c \"echo -1 > /proc/sys/kernel/perf_event_paranoid\"'\n"
				    "Hint:\tThe current value is %d.", value);
1273 1274 1275 1276 1277 1278 1279 1280
		break;
	default:
		scnprintf(buf, size, "%s", emsg);
		break;
	}

	return 0;
}
1281 1282 1283 1284 1285 1286 1287 1288 1289 1290

void perf_evlist__to_front(struct perf_evlist *evlist,
			   struct perf_evsel *move_evsel)
{
	struct perf_evsel *evsel, *n;
	LIST_HEAD(move);

	if (move_evsel == perf_evlist__first(evlist))
		return;

1291
	evlist__for_each_safe(evlist, n, evsel) {
1292 1293 1294 1295 1296 1297
		if (evsel->leader == move_evsel->leader)
			list_move_tail(&evsel->node, &move);
	}

	list_splice(&move, &evlist->entries);
}