staticvoid record__read_auxtrace_snapshot(struct record *rec, bool on_exit)
{
pr_debug("Recording AUX area tracing snapshot\n"); if (record__auxtrace_read_snapshot_all(rec) < 0) {
trigger_error(&auxtrace_snapshot_trigger);
} else { if (auxtrace_record__snapshot_finish(rec->itr, on_exit))
trigger_error(&auxtrace_snapshot_trigger); else
trigger_ready(&auxtrace_snapshot_trigger);
}
}
staticint record__auxtrace_snapshot_exit(struct record *rec)
{ if (trigger_is_error(&auxtrace_snapshot_trigger)) return0;
if (!auxtrace_record__snapshot_started &&
auxtrace_record__snapshot_start(rec->itr)) return -1;
record__read_auxtrace_snapshot(rec, true); if (trigger_is_error(&auxtrace_snapshot_trigger)) return -1;
return0;
}
staticint record__auxtrace_init(struct record *rec)
{ int err;
if ((rec->opts.auxtrace_snapshot_opts || rec->opts.auxtrace_sample_opts)
&& record__threads_enabled(rec)) {
pr_err("AUX area tracing options are not available in parallel streaming mode.\n"); return -EINVAL;
}
if (!rec->itr) {
rec->itr = auxtrace_record__init(rec->evlist, &err); if (err) return err;
}
err = auxtrace_parse_snapshot_options(rec->itr, &rec->opts,
rec->opts.auxtrace_snapshot_opts); if (err) return err;
err = auxtrace_parse_sample_options(rec->itr, rec->evlist, &rec->opts,
rec->opts.auxtrace_sample_opts); if (err) return err;
err = auxtrace_parse_aux_action(rec->evlist); if (err) return err;
return auxtrace_parse_filters(rec->evlist);
}
#else
staticinline int record__auxtrace_mmap_read(struct record *rec __maybe_unused, struct mmap *map __maybe_unused)
{ return0;
}
staticvoid record__free_thread_data(struct record *rec)
{ int t; struct record_thread *thread_data = rec->thread_data;
if (thread_data == NULL) return;
for (t = 0; t < rec->nr_threads; t++) {
record__thread_data_close_pipes(&thread_data[t]);
zfree(&thread_data[t].maps);
zfree(&thread_data[t].overwrite_maps);
fdarray__exit(&thread_data[t].pollfd);
}
zfree(&rec->thread_data);
}
staticint record__map_thread_evlist_pollfd_indexes(struct record *rec, int evlist_pollfd_index, int thread_pollfd_index)
{
size_t x = rec->index_map_cnt;
for (i = 0; i < rec->index_map_cnt; i++) { int e_pos = rec->index_map[i].evlist_pollfd_index; int t_pos = rec->index_map[i].thread_pollfd_index;
if (e_entries[e_pos].fd != t_entries[t_pos].fd ||
e_entries[e_pos].events != t_entries[t_pos].events) {
pr_err("Thread and evlist pollfd index mismatch\n");
err = -EINVAL; continue;
}
e_entries[e_pos].revents = t_entries[t_pos].revents;
} return err;
}
staticint record__dup_non_perf_events(struct record *rec, struct evlist *evlist, struct record_thread *thread_data)
{ struct fdarray *fda = &evlist->core.pollfd; int i, ret;
for (i = 0; i < fda->nr; i++) { if (!(fda->priv[i].flags & fdarray_flag__non_perf_event)) continue;
ret = fdarray__dup_entry_from(&thread_data->pollfd, i, fda); if (ret < 0) {
pr_err("Failed to duplicate descriptor in main thread pollfd\n"); return ret;
}
pr_debug2("thread_data[%p]: pollfd[%d] <- non_perf_event fd=%d\n",
thread_data, ret, fda->entries[i].fd);
ret = record__map_thread_evlist_pollfd_indexes(rec, i, ret); if (ret < 0) {
pr_err("Failed to map thread and evlist pollfd indexes\n"); return ret;
}
} return0;
}
staticint record__alloc_thread_data(struct record *rec, struct evlist *evlist)
{ int t, ret; struct record_thread *thread_data;
rec->thread_data = zalloc(rec->nr_threads * sizeof(*(rec->thread_data))); if (!rec->thread_data) {
pr_err("Failed to allocate thread data\n"); return -ENOMEM;
}
thread_data = rec->thread_data;
for (t = 0; t < rec->nr_threads; t++)
record__thread_data_init_pipes(&thread_data[t]);
for (t = 0; t < rec->nr_threads; t++) {
thread_data[t].rec = rec;
thread_data[t].mask = &rec->thread_masks[t];
ret = record__thread_data_init_maps(&thread_data[t], evlist); if (ret) {
pr_err("Failed to initialize thread[%d] maps\n", t); goto out_free;
}
ret = record__thread_data_init_pollfd(&thread_data[t], evlist); if (ret) {
pr_err("Failed to initialize thread[%d] pollfd\n", t); goto out_free;
} if (t) {
thread_data[t].tid = -1;
ret = record__thread_data_open_pipes(&thread_data[t]); if (ret) {
pr_err("Failed to open thread[%d] communication pipes\n", t); goto out_free;
}
ret = fdarray__add(&thread_data[t].pollfd, thread_data[t].pipes.msg[0],
POLLIN | POLLERR | POLLHUP, fdarray_flag__nonfilterable); if (ret < 0) {
pr_err("Failed to add descriptor to thread[%d] pollfd\n", t); goto out_free;
}
thread_data[t].ctlfd_pos = ret;
pr_debug2("thread_data[%p]: pollfd[%d] <- ctl_fd=%d\n",
thread_data, thread_data[t].ctlfd_pos,
thread_data[t].pipes.msg[0]);
} else {
thread_data[t].tid = gettid();
ret = record__dup_non_perf_events(rec, evlist, &thread_data[t]); if (ret < 0) goto out_free;
thread_data[t].ctlfd_pos = -1; /* Not used */
}
}
return0;
out_free:
record__free_thread_data(rec);
return ret;
}
staticint record__mmap_evlist(struct record *rec, struct evlist *evlist)
{ int i, ret; struct record_opts *opts = &rec->opts; bool auxtrace_overwrite = opts->auxtrace_snapshot_mode ||
opts->auxtrace_sample_mode; char msg[512];
if (opts->affinity != PERF_AFFINITY_SYS)
cpu__setup_cpunode_map();
if (evlist__mmap_ex(evlist, opts->mmap_pages,
opts->auxtrace_mmap_pages,
auxtrace_overwrite,
opts->nr_cblocks, opts->affinity,
opts->mmap_flush, opts->comp_level) < 0) { if (errno == EPERM) {
pr_err("Permission error mapping pages.\n" "Consider increasing " "/proc/sys/kernel/perf_event_mlock_kb,\n" "or try again with a smaller value of -m/--mmap_pages.\n" "(current value: %u,%u)\n",
opts->mmap_pages, opts->auxtrace_mmap_pages); return -errno;
} else {
pr_err("failed to mmap with %d (%s)\n", errno,
str_error_r(errno, msg, sizeof(msg))); if (errno) return -errno; else return -EINVAL;
}
}
if (evlist__initialize_ctlfd(evlist, opts->ctl_fd, opts->ctl_fd_ack)) return -1;
ret = record__alloc_thread_data(rec, evlist); if (ret) return ret;
if (record__threads_enabled(rec)) {
ret = perf_data__create_dir(&rec->data, evlist->core.nr_mmaps); if (ret) {
pr_err("Failed to create data directory: %s\n", strerror(-ret)); return ret;
} for (i = 0; i < evlist->core.nr_mmaps; i++) { if (evlist->mmap)
evlist->mmap[i].file = &rec->data.dir.files[i]; if (evlist->overwrite_mmap)
evlist->overwrite_mmap[i].file = &rec->data.dir.files[i];
}
}
return0;
}
staticint record__mmap(struct record *rec)
{ return record__mmap_evlist(rec, rec->evlist);
}
if (symbol_conf.kptr_restrict && !evlist__exclude_kernel(evlist)) {
pr_warning( "WARNING: Kernel address maps (/proc/{kallsyms,modules}) are restricted,\n" "check /proc/sys/kernel/kptr_restrict and /proc/sys/kernel/perf_event_paranoid.\n\n" "Samples in kernel functions may not be resolved if a suitable vmlinux\n" "file is not found in the buildid cache or in the vmlinux path.\n\n" "Samples in kernel modules won't be resolved at all.\n\n" "If some relocation was applied (e.g. kexec) symbols may be misresolved\n" "even with a suitable vmlinux or kallsyms file.\n\n");
}
if (evlist__apply_filters(evlist, &pos, &opts->target)) {
pr_err("failed to set filter \"%s\" on event %s with %d (%s)\n",
pos->filter ?: "BPF", evsel__name(pos), errno,
str_error_r(errno, msg, sizeof(msg)));
rc = -1; goto out;
}
staticint record__synthesize(struct record *rec, bool tail);
staticint
record__switch_output(struct record *rec, bool at_exit)
{ struct perf_data *data = &rec->data; char *new_filename = NULL; int fd, err;
/* Same Size: "2015122520103046"*/ char timestamp[] = "InvalidTimestamp";
record__aio_mmap_read_sync(rec);
write_finished_init(rec, true);
record__synthesize(rec, true); if (target__none(&rec->opts.target))
record__synthesize_workload(rec, true);
rec->samples = 0;
record__finish_output(rec);
err = fetch_current_timestamp(timestamp, sizeof(timestamp)); if (err) {
pr_err("Failed to get current timestamp\n"); return -EINVAL;
}
if (data->is_pipe) {
err = perf_event__synthesize_for_pipe(tool, session, data,
process_synthesized_event); if (err < 0) goto out;
rec->bytes_written += err;
}
err = perf_event__synth_time_conv(record__pick_pc(rec), tool,
process_synthesized_event, machine); if (err) goto out;
/* Synthesize id_index before auxtrace_info */
err = perf_event__synthesize_id_index(tool,
process_synthesized_event,
session->evlist, machine); if (err) goto out;
if (rec->opts.full_auxtrace) {
err = perf_event__synthesize_auxtrace_info(rec->itr, tool,
session, process_synthesized_event); if (err) goto out;
}
if (!evlist__exclude_kernel(rec->evlist)) {
err = perf_event__synthesize_kernel_mmap(tool, process_synthesized_event,
machine);
WARN_ONCE(err < 0, "Couldn't record kernel reference relocation symbol\n" "Symbol resolution may be skewed if relocation was used (e.g. kexec).\n" "Check /proc/kallsyms permission or run as root.\n");
err = perf_event__synthesize_modules(tool, process_synthesized_event,
machine);
WARN_ONCE(err < 0, "Couldn't record kernel module information.\n" "Symbol resolution may be skewed if relocation was used (e.g. kexec).\n" "Check /proc/modules permission or run as root.\n");
}
if (perf_guest) {
machines__process_guests(&session->machines,
perf_event__synthesize_guest_os, tool);
}
err = perf_event__synthesize_extra_attr(&rec->tool,
rec->evlist,
process_synthesized_event,
data->is_pipe); if (err) goto out;
staticint record__setup_sb_evlist(struct record *rec)
{ struct record_opts *opts = &rec->opts;
if (rec->sb_evlist != NULL) { /* *Wegethereif--switch-output-eventpopulatedthe *sb_evlist,soassociateacallbackthatwillsendaSIGUSR2 *tothemainthread.
*/
evlist__set_cb(rec->sb_evlist, record__process_signal_event, rec);
rec->thread_id = pthread_self();
} #ifdef HAVE_LIBBPF_SUPPORT if (!opts->no_bpf_event) { if (rec->sb_evlist == NULL) {
rec->sb_evlist = evlist__new();
if (rec->sb_evlist == NULL) {
pr_err("Couldn't create side band evlist.\n."); return -1;
}
}
if (evlist__add_bpf_sb_event(rec->sb_evlist, perf_session__env(rec->session))) {
pr_err("Couldn't ask for PERF_RECORD_BPF_EVENT side band events.\n."); return -1;
}
} #endif if (evlist__start_sb_thread(rec->sb_evlist, &rec->opts.target)) {
pr_debug("Couldn't start the BPF side band thread:\nBPF programs starting from now on won't be annotatable\n");
opts->no_bpf_event = true;
}
if (rec->timestamp_filename && perf_data__is_pipe(data)) {
rec->timestamp_filename = false;
pr_warning("WARNING: --timestamp-filename option is not available in pipe mode.\n");
}
/* Debug message used by test scripts */
pr_debug3("perf record opening and mmapping events\n"); if (record__open(rec) != 0) {
err = -1; goto out_free_threads;
} /* Debug message used by test scripts */
pr_debug3("perf record done opening and mmapping events\n");
env->comp_mmap_len = session->evlist->core.mmap_len;
if (rec->opts.kcore) {
err = record__kcore_copy(&session->machines.host, data); if (err) {
pr_err("ERROR: Failed to copy kcore\n"); goto out_free_threads;
}
}
/* *Normallyperf_session__newwoulddothis,butitdoesn'thavethe *evlist.
*/ if (rec->tool.ordered_events && !evlist__sample_id_all(rec->evlist)) {
pr_warning("WARNING: No sample_id_all support, falling back to unordered processing\n");
rec->tool.ordered_events = false;
}
if (evlist__nr_groups(rec->evlist) == 0)
perf_header__clear_feat(&session->header, HEADER_GROUP_DESC);
if (data->is_pipe) {
err = perf_header__write_pipe(fd); if (err < 0) goto out_free_threads;
} else {
err = perf_session__write_header(session, rec->evlist, fd, false); if (err < 0) goto out_free_threads;
}
err = record__update_evlist_pollfd_from_thread(rec, rec->evlist, thread); if (err) goto out_child;
}
if (evlist__ctlfd_process(rec->evlist, &cmd) > 0) { switch (cmd) { case EVLIST_CTL_CMD_SNAPSHOT:
hit_auxtrace_snapshot_trigger(rec);
evlist__ctlfd_ack(rec->evlist); break; case EVLIST_CTL_CMD_STOP:
done = 1; break; case EVLIST_CTL_CMD_ACK: case EVLIST_CTL_CMD_UNSUPPORTED: case EVLIST_CTL_CMD_ENABLE: case EVLIST_CTL_CMD_DISABLE: case EVLIST_CTL_CMD_EVLIST: case EVLIST_CTL_CMD_PING: default: break;
}
}
err = event_enable_timer__process(rec->evlist->eet); if (err < 0) goto out_child; if (err) {
err = 0;
done = 1;
}
if (rec->session->bytes_transferred && rec->session->bytes_compressed) {
ratio = (float)rec->session->bytes_transferred/(float)rec->session->bytes_compressed;
env->comp_ratio = ratio + 0.5;
}
if (forks) { int exit_status;
if (!child_finished)
kill(rec->evlist->workload.pid, SIGTERM);
wait(&exit_status);
if (err < 0)
status = err; elseif (WIFEXITED(exit_status))
status = WEXITSTATUS(exit_status); elseif (WIFSIGNALED(exit_status))
signr = WTERMSIG(exit_status);
} else
status = err;
if (rec->off_cpu)
rec->bytes_written += off_cpu_write(rec->session);
record__read_lost_samples(rec);
record__synthesize(rec, true); /* this will be recalculated during process_buildids() */
rec->samples = 0;
if (!err) { if (!rec->timestamp_filename) {
record__finish_output(rec);
} else {
fd = record__switch_output(rec, true); if (fd < 0) {
status = fd; goto out_delete_session;
}
}
}
ret = parse_callchain_record_opt(arg, callchain); if (!ret) { /* Enable data address sampling for DWARF unwind. */ if (callchain->record_mode == CALLCHAIN_DWARF)
record->sample_address = true;
callchain_debug(callchain);
}
return ret;
}
int record_parse_callchain_opt(conststruct option *opt, constchar *arg, int unset)
{ return record_opts__parse_callchain(opt->value, &callchain_param, arg, unset);
}
int record_callchain_opt(conststruct option *opt, constchar *arg __maybe_unused, int unset __maybe_unused)
{ struct callchain_param *callchain = opt->value;
callchain->enabled = true;
if (callchain->record_mode == CALLCHAIN_NONE)
callchain->record_mode = CALLCHAIN_FP;
/* *Ifwe'reusing--switch-output-events,thenweimplyits *--switch-output=signal,aswe'llsendaSIGUSR2fromthesideband *threadtoitsparent.
*/ if (rec->switch_output_event_set) { if (record__threads_enabled(rec)) {
pr_warning("WARNING: --switch-output-event option is not available in parallel streaming mode.\n"); return0;
} goto do_signal;
}
if (!s->set) return0;
if (record__threads_enabled(rec)) {
pr_warning("WARNING: --switch-output option is not available in parallel streaming mode.\n"); return0;
}
if (!strcmp(s->str, "signal")) {
do_signal:
s->signal = true;
pr_debug("switch-output with SIGUSR2 signal\n"); goto enabled;
}
val = parse_tag_value(s->str, tags_size); if (val != (unsignedlong) -1) {
s->size = val;
pr_debug("switch-output with %s size threshold\n", s->str); goto enabled;
}
val = parse_tag_value(s->str, tags_time); if (val != (unsignedlong) -1) {
s->time = val;
pr_debug("switch-output with %s time threshold (%lu seconds)\n",
s->str, s->time); goto enabled;
}
/* *XXXWillstayaglobalvariabletillwefixbuiltin-script.ctostopmessing *withitandswitchtousethelibraryfunctionsinperf_evlistthatcame *frombuiltin-record.c,i.e.userecord_opts, *evlist__prepare_workload,etcinsteadoffork+exec'in'perfrecord', *usingpipes,etc.
*/ staticstruct option __record_options[] = {
OPT_CALLBACK('e', "event", &parse_events_option_args, "event", "event selector. use 'perf list' to list available events",
parse_events_option),
OPT_CALLBACK(0, "filter", &record.evlist, "filter", "event filter", parse_filter),
OPT_BOOLEAN(0, "latency", &record.latency, "Enable data collection for latency profiling.\n" "\t\t\t Use perf report --latency for latency-centric profile."),
OPT_CALLBACK_NOOPT(0, "exclude-perf", &record.evlist,
NULL, "don't record events from perf itself",
exclude_perf),
OPT_STRING('p', "pid", &record.opts.target.pid, "pid", "record events on existing process id"),
OPT_STRING('t', "tid", &record.opts.target.tid, "tid", "record events on existing thread id"),
OPT_INTEGER('r', "realtime", &record.realtime_prio, "collect data with this RT SCHED_FIFO priority"),
OPT_BOOLEAN(0, "no-buffering", &record.opts.no_buffering, "collect data without buffering"),
OPT_BOOLEAN('R', "raw-samples", &record.opts.raw_samples, "collect raw sample records from all opened counters"),
OPT_BOOLEAN('a', "all-cpus", &record.opts.target.system_wide, "system-wide collection from all CPUs"),
OPT_STRING('C', "cpu", &record.opts.target.cpu_list, "cpu", "list of cpus to monitor"),
OPT_U64('c', "count", &record.opts.user_interval, "event period to sample"),
OPT_STRING('o', "output", &record.data.path, "file", "output file name"),
OPT_BOOLEAN_SET('i', "no-inherit", &record.opts.no_inherit,
&record.opts.no_inherit_set, "child tasks do not inherit counters"),
OPT_BOOLEAN(0, "tail-synthesize", &record.opts.tail_synthesize, "synthesize non-sample events at the end of output"),
OPT_BOOLEAN(0, "overwrite", &record.opts.overwrite, "use overwrite mode"),
OPT_BOOLEAN(0, "no-bpf-event", &record.opts.no_bpf_event, "do not record bpf events"),
OPT_BOOLEAN(0, "strict-freq", &record.opts.strict_freq, "Fail if the specified frequency can't be used"),
OPT_CALLBACK('F', "freq", &record.opts, "freq or 'max'", "profile at this frequency",
record__parse_freq),
OPT_CALLBACK('m', "mmap-pages", &record.opts, "pages[,pages]", "number of mmap data pages and AUX area tracing mmap pages",
record__parse_mmap_pages),
OPT_CALLBACK(0, "mmap-flush", &record.opts, "number", "Minimal number of bytes that is extracted from mmap data pages (default: 1)",
record__mmap_flush_parse),
OPT_CALLBACK_NOOPT('g', NULL, &callchain_param,
NULL, "enables call-graph recording" ,
&record_callchain_opt),
OPT_CALLBACK(0, "call-graph", &record.opts, "record_mode[,record_size]", record_callchain_help,
&record_parse_callchain_opt),
OPT_INCR('v', "verbose", &verbose, "be more verbose (show counter open errors, etc)"),
OPT_BOOLEAN('q', "quiet", &quiet, "don't print any warnings or messages"),
OPT_BOOLEAN('s', "stat", &record.opts.inherit_stat, "per thread counts"),
OPT_BOOLEAN('d', "data", &record.opts.sample_address, "Record the sample addresses"),
OPT_BOOLEAN(0, "phys-data", &record.opts.sample_phys_addr, "Record the sample physical addresses"),
OPT_BOOLEAN(0, "data-page-size", &record.opts.sample_data_page_size, "Record the sampled data address data page size"),
OPT_BOOLEAN(0, "code-page-size", &record.opts.sample_code_page_size, "Record the sampled code address (ip) page size"),
OPT_BOOLEAN(0, "sample-mem-info", &record.opts.sample_data_src, "Record the data source for memory operations"),
OPT_BOOLEAN(0, "sample-cpu", &record.opts.sample_cpu, "Record the sample cpu"),
OPT_BOOLEAN(0, "sample-identifier", &record.opts.sample_identifier, "Record the sample identifier"),
OPT_BOOLEAN_SET('T', "timestamp", &record.opts.sample_time,
&record.opts.sample_time_set, "Record the sample timestamps"),
OPT_BOOLEAN_SET('P', "period", &record.opts.period, &record.opts.period_set, "Record the sample period"),
OPT_BOOLEAN('n', "no-samples", &record.opts.no_samples, "don't sample"),
OPT_BOOLEAN_SET('N', "no-buildid-cache", &record.no_buildid_cache,
&record.no_buildid_cache_set, "do not update the buildid cache"),
OPT_BOOLEAN_SET('B', "no-buildid", &record.no_buildid,
&record.no_buildid_set, "do not collect buildids in perf.data"),
OPT_CALLBACK('G', "cgroup", &record.evlist, "name", "monitor event in cgroup name only",
parse_cgroups),
OPT_CALLBACK('D', "delay", &record, "ms", "ms to wait before starting measurement after program start (-1: start with events disabled), " "or ranges of time to enable events e.g. '-D 10-20,30-40'",
record__parse_event_enable_time),
OPT_BOOLEAN(0, "kcore", &record.opts.kcore, "copy /proc/kcore"),
OPT_STRING('u', "uid", &record.uid_str, "user", "user to profile"),
OPT_CALLBACK_NOOPT('b', "branch-any", &record.opts.branch_stack, "branch any", "sample any taken branches",
parse_branch_stack),
OPT_CALLBACK('j', "branch-filter", &record.opts.branch_stack, "branch filter mask", "branch stack filter modes",
parse_branch_stack),
OPT_BOOLEAN('W', "weight", &record.opts.sample_weight, "sample by weight (on special events only)"),
OPT_BOOLEAN(0, "transaction", &record.opts.sample_transaction, "sample transaction flags (special events only)"),
OPT_BOOLEAN(0, "per-thread", &record.opts.target.per_thread, "use per-thread mmaps"),
OPT_CALLBACK_OPTARG('I', "intr-regs", &record.opts.sample_intr_regs, NULL, "any register", "sample selected machine registers on interrupt," " use '-I?' to list register names", parse_intr_regs),
OPT_CALLBACK_OPTARG(0, "user-regs", &record.opts.sample_user_regs, NULL, "any register", "sample selected machine registers in user space," " use '--user-regs=?' to list register names", parse_user_regs),
OPT_BOOLEAN(0, "running-time", &record.opts.running_time, "Record running/enabled time of read (:S) events"),
OPT_CALLBACK('k', "clockid", &record.opts, "clockid", "clockid to use for events, see clock_gettime()",
parse_clockid),
OPT_STRING_OPTARG('S', "snapshot", &record.opts.auxtrace_snapshot_opts, "opts", "AUX area tracing Snapshot Mode", ""),
OPT_STRING_OPTARG(0, "aux-sample", &record.opts.auxtrace_sample_opts, "opts", "sample AUX area", ""),
OPT_UINTEGER(0, "proc-map-timeout", &proc_map_timeout, "per thread proc mmap processing timeout in ms"),
OPT_BOOLEAN(0, "namespaces", &record.opts.record_namespaces, "Record namespaces events"),
OPT_BOOLEAN(0, "all-cgroups", &record.opts.record_cgroup, "Record cgroup events"),
OPT_BOOLEAN_SET(0, "switch-events", &record.opts.record_switch_events,
&record.opts.record_switch_events_set, "Record context switch events"),
OPT_BOOLEAN_FLAG(0, "all-kernel", &record.opts.all_kernel, "Configure all used events to run in kernel space.",
PARSE_OPT_EXCLUSIVE),
OPT_BOOLEAN_FLAG(0, "all-user", &record.opts.all_user, "Configure all used events to run in user space.",
PARSE_OPT_EXCLUSIVE),
OPT_BOOLEAN(0, "kernel-callchains", &record.opts.kernel_callchains, "collect kernel callchains"),
OPT_BOOLEAN(0, "user-callchains", &record.opts.user_callchains, "collect user callchains"),
OPT_STRING(0, "vmlinux", &symbol_conf.vmlinux_name, "file", "vmlinux pathname"),
OPT_BOOLEAN(0, "buildid-all", &record.buildid_all, "Record build-id of all DSOs regardless of hits"),
OPT_BOOLEAN_SET(0, "buildid-mmap", &record.buildid_mmap, &record.buildid_mmap_set, "Record build-id in mmap events and skip build-id processing."),
OPT_BOOLEAN(0, "timestamp-filename", &record.timestamp_filename, "append timestamp to output filename"),
OPT_BOOLEAN(0, "timestamp-boundary", &record.timestamp_boundary, "Record timestamp boundary (time of first/last samples)"),
OPT_STRING_OPTARG_SET(0, "switch-output", &record.switch_output.str,
&record.switch_output.set, "signal or size[BKMG] or time[smhd]", "Switch output when receiving SIGUSR2 (signal) or cross a size or time threshold", "signal"),
OPT_CALLBACK_SET(0, "switch-output-event", &switch_output_parse_events_option_args,
&record.switch_output_event_set, "switch output event", "switch output event selector. use 'perf list' to list available events",
parse_events_option_new_evlist),
OPT_INTEGER(0, "switch-max-files", &record.switch_output.num_files, "Limit number of switch output generated files"),
OPT_BOOLEAN(0, "dry-run", &dry_run, "Parse options then exit"), #ifdef HAVE_AIO_SUPPORT
OPT_CALLBACK_OPTARG(0, "aio", &record.opts,
&nr_cblocks_default, "n", "Use <n> control blocks in asynchronous trace writing mode (default: 1, max: 4)",
record__aio_parse), #endif
OPT_CALLBACK(0, "affinity", &record.opts, "node|cpu", "Set affinity mask of trace reading thread to NUMA node cpu mask or cpu of processed mmap buffer",
record__parse_affinity), #ifdef HAVE_ZSTD_SUPPORT
OPT_CALLBACK_OPTARG('z', "compression-level", &record.opts, &comp_level_default, "n", "Compress records using specified level (default: 1 - fastest compression, 22 - greatest compression)",
record__parse_comp_level), #endif
OPT_CALLBACK(0, "max-size", &record.output_max_size, "size", "Limit the maximum size of the output file", parse_output_max_size),
OPT_UINTEGER(0, "num-thread-synthesize",
&record.opts.nr_threads_synthesize, "number of threads to run for event synthesis"), #ifdef HAVE_LIBPFM
OPT_CALLBACK(0, "pfm-events", &record.evlist, "event", "libpfm4 event selector. use 'perf list' to list available events",
parse_libpfm_events_option), #endif
OPT_CALLBACK(0, "control", &record.opts, "fd:ctl-fd[,ack-fd] or fifo:ctl-fifo[,ack-fifo]", "Listen on ctl-fd descriptor for command to control measurement ('enable': enable events, 'disable': disable events,\n" "\t\t\t 'snapshot': AUX area tracing snapshot).\n" "\t\t\t Optionally send control command completion ('ack\\n') to ack-fd descriptor.\n" "\t\t\t Alternatively, ctl-fifo / ack-fifo will be opened and used as ctl-fd / ack-fd.",
parse_control_option),
OPT_CALLBACK(0, "synth", &record.opts, "no|all|task|mmap|cgroup", "Fine-tune event synthesis: default=all", parse_record_synth_option),
OPT_STRING_OPTARG_SET(0, "debuginfod", &record.debuginfod.urls,
&record.debuginfod.set, "debuginfod urls", "Enable debuginfod data retrieval from DEBUGINFOD_URLS or specified urls", "system"),
OPT_CALLBACK_OPTARG(0, "threads", &record.opts, NULL, "spec", "write collected trace data into several data files using parallel threads",
record__parse_threads),
OPT_BOOLEAN(0, "off-cpu", &record.off_cpu, "Enable off-cpu analysis"),
OPT_STRING(0, "setup-filter", &record.filter_action, "pin|unpin", "BPF filter action"),
OPT_CALLBACK(0, "off-cpu-thresh", &record.opts, "ms", "Dump off-cpu samples if off-cpu time exceeds this threshold (in milliseconds). (Default: 500ms)",
record__parse_off_cpu_thresh),
OPT_END()
};
struct option *record_options = __record_options;
static int record__mmap_cpu_mask_init(struct mmap_cpu_mask *mask, struct perf_cpu_map *cpus)
{
struct perf_cpu cpu;
int idx;
if (cpu_map__is_dummy(cpus))
return 0;
perf_cpu_map__for_each_cpu_skip_any(cpu, idx, cpus) {
/* Return ENODEV is input cpu is greater than max cpu */
if ((unsigned long)cpu.cpu > mask->nbits)
return -ENODEV;
__set_bit(cpu.cpu, mask->bits);
}
static int record__init_thread_masks_spec(struct record *rec, struct perf_cpu_map *cpus,
const char **maps_spec, const char **affinity_spec,
u32 nr_spec)
{
u32 s;
int ret = 0, t = 0;
struct mmap_cpu_mask cpus_mask;
struct thread_mask thread_mask, full_mask, *thread_masks;
ret = record__mmap_cpu_mask_alloc(&cpus_mask, cpu__max_cpu().cpu);
if (ret) {
pr_err("Failed to allocate CPUs mask\n");
return ret;
}
ret = record__mmap_cpu_mask_init(&cpus_mask, cpus);
if (ret) {
pr_err("Failed to init cpu mask\n");
goto out_free_cpu_mask;
}
ret = record__thread_mask_alloc(&full_mask, cpu__max_cpu().cpu);
if (ret) {
pr_err("Failed to allocate full mask\n");
goto out_free_cpu_mask;
}
ret = record__thread_mask_alloc(&thread_mask, cpu__max_cpu().cpu);
if (ret) {
pr_err("Failed to allocate thread mask\n");
goto out_free_full_and_cpu_masks;
}
for (s = 0; s < nr_spec; s++) {
ret = record__mmap_cpu_mask_init_spec(&thread_mask.maps, maps_spec[s]);
if (ret) {
pr_err("Failed to initialize maps thread mask\n");
goto out_free;
}
ret = record__mmap_cpu_mask_init_spec(&thread_mask.affinity, affinity_spec[s]);
if (ret) {
pr_err("Failed to initialize affinity thread mask\n");
goto out_free;
}
/* ignore invalid CPUs but do not allow empty masks */
if (!bitmap_and(thread_mask.maps.bits, thread_mask.maps.bits,
cpus_mask.bits, thread_mask.maps.nbits)) {
pr_err("Empty maps mask: %s\n", maps_spec[s]);
ret = -EINVAL;
goto out_free;
}
if (!bitmap_and(thread_mask.affinity.bits, thread_mask.affinity.bits,
cpus_mask.bits, thread_mask.affinity.nbits)) {
pr_err("Empty affinity mask: %s\n", affinity_spec[s]);
ret = -EINVAL;
goto out_free;
}
/* do not allow intersection with other masks (full_mask) */
if (bitmap_intersects(thread_mask.maps.bits, full_mask.maps.bits,
thread_mask.maps.nbits)) {
pr_err("Intersecting maps mask: %s\n", maps_spec[s]);
ret = -EINVAL;
goto out_free;
}
if (bitmap_intersects(thread_mask.affinity.bits, full_mask.affinity.bits,
thread_mask.affinity.nbits)) {
pr_err("Intersecting affinity mask: %s\n", affinity_spec[s]);
ret = -EINVAL;
goto out_free;
}
static int record__init_thread_core_masks(struct record *rec, struct perf_cpu_map *cpus)
{
int ret;
struct cpu_topology *topo;
topo = cpu_topology__new();
if (!topo) {
pr_err("Failed to allocate CPU topology\n");
return -ENOMEM;
}
ret = record__init_thread_masks_spec(rec, cpus, topo->core_cpus_list,
topo->core_cpus_list, topo->core_cpus_lists);
cpu_topology__delete(topo);
return ret;
}
static int record__init_thread_package_masks(struct record *rec, struct perf_cpu_map *cpus)
{
int ret;
struct cpu_topology *topo;
topo = cpu_topology__new();
if (!topo) {
pr_err("Failed to allocate CPU topology\n");
return -ENOMEM;
}
ret = record__init_thread_masks_spec(rec, cpus, topo->package_cpus_list,
topo->package_cpus_list, topo->package_cpus_lists);
cpu_topology__delete(topo);
return ret;
}
static int record__init_thread_numa_masks(struct record *rec, struct perf_cpu_map *cpus)
{
u32 s;
int ret;
const char **spec;
struct numa_topology *topo;
topo = numa_topology__new();
if (!topo) {
pr_err("Failed to allocate NUMA topology\n");
return -ENOMEM;
}
spec = zalloc(topo->nr * sizeof(char *));
if (!spec) {
pr_err("Failed to allocate NUMA spec\n");
ret = -ENOMEM;
goto out_delete_topo;
}
for (s = 0; s < topo->nr; s++)
spec[s] = topo->nodes[s].cpus;
ret = record__init_thread_masks_spec(rec, cpus, spec, spec, topo->nr);
out_free:
free(dup_mask);
for (s = 0; s < nr_spec; s++) {
if (maps_spec)
free(maps_spec[s]);
if (affinity_spec)
free(affinity_spec[s]);
}
free(affinity_spec);
free(maps_spec);
return ret;
}
static int record__init_thread_default_masks(struct record *rec, struct perf_cpu_map *cpus)
{
int ret;
ret = record__alloc_thread_masks(rec, 1, cpu__max_cpu().cpu);
if (ret)
return ret;
if (record__mmap_cpu_mask_init(&rec->thread_masks->maps, cpus))
return -ENODEV;
rec->nr_threads = 1;
return 0;
}
static int record__init_thread_masks(struct record *rec)
{
int ret = 0;
struct perf_cpu_map *cpus = rec->evlist->core.all_cpus;
if (!record__threads_enabled(rec))
return record__init_thread_default_masks(rec, cpus);
if (evlist__per_thread(rec->evlist)) {
pr_err("--per-thread option is mutually exclusive to parallel streaming mode.\n");
return -EINVAL;
}
switch (rec->opts.threads_spec) {
case THREAD_SPEC__CPU:
ret = record__init_thread_cpu_masks(rec, cpus);
break;
case THREAD_SPEC__CORE:
ret = record__init_thread_core_masks(rec, cpus);
break;
case THREAD_SPEC__PACKAGE:
ret = record__init_thread_package_masks(rec, cpus);
break;
case THREAD_SPEC__NUMA:
ret = record__init_thread_numa_masks(rec, cpus);
break;
case THREAD_SPEC__USER:
ret = record__init_thread_user_masks(rec, cpus);
break;
default:
break;
}
return ret;
}
int cmd_record(int argc, const char **argv)
{
int err;
struct record *rec = &record;
char errbuf[BUFSIZ];
setlocale(LC_ALL, "");
#ifndef HAVE_BPF_SKEL
# define set_nobuild(s, l, m, c) set_option_nobuild(record_options, s, l, m, c)
set_nobuild('\0', "off-cpu", "no BUILD_BPF_SKEL=1", true);
# undef set_nobuild
#endif
/* Disable eager loading of kernel symbols that adds overhead to perf record. */
symbol_conf.lazy_load_kernel_maps = true;
rec->opts.affinity = PERF_AFFINITY_SYS;
rec->evlist = evlist__new();
if (rec->evlist == NULL)
return -ENOMEM;
err = perf_config(perf_record_config, rec);
if (err)
return err;
argc = parse_options(argc, argv, record_options, record_usage,
PARSE_OPT_STOP_AT_NON_OPTION);
if (quiet)
perf_quiet_option();
err = symbol__validate_sym_arguments();
if (err)
return err;
perf_debuginfod_setup(&record.debuginfod);
/* Make system wide (-a) the default target. */
if (!argc && target__none(&rec->opts.target))
rec->opts.target.system_wide = true;
if (nr_cgroups && !rec->opts.target.system_wide) {
usage_with_options_msg(record_usage, record_options,
"cgroup monitoring only available in system-wide mode");
}
if (record.latency) {
/*
* There is no fundamental reason why latency profiling
* can't work for system-wide mode, but exact semantics
* and details are to be defined.
* See the following thread for details:
* https://lore.kernel.org/all/Z4XDJyvjiie3howF@google.com/
*/
if (record.opts.target.system_wide) {
pr_err("Failed: latency profiling is not supported with system-wide collection.\n");
err = -EINVAL;
goto out_opts;
}
record.opts.record_switch_events = true;
}
if (!rec->buildid_mmap) {
pr_debug("Disabling build id in synthesized mmap2 events.\n");
symbol_conf.no_buildid_mmap2 = true;
} else if (rec->buildid_mmap_set) {
/*
* Explicitly passing --buildid-mmap disables buildid processing
* and cache generation.
*/
rec->no_buildid = true;
}
if (rec->buildid_mmap && !perf_can_record_build_id()) {
pr_warning("Missing support for build id in kernel mmap events.\n"
"Disable this warning with --no-buildid-mmap\n");
rec->buildid_mmap = false;
}
if (rec->buildid_mmap) {
/* Enable perf_event_attr::build_id bit. */
rec->opts.build_id = true;
}
if (rec->opts.record_cgroup && !perf_can_record_cgroup()) {
pr_err("Kernel has no cgroup sampling support.\n");
err = -EINVAL;
goto out_opts;
}
if (rec->opts.kcore)
rec->opts.text_poke = true;
if (rec->opts.kcore || record__threads_enabled(rec))
rec->data.is_dir = true;
if (record__threads_enabled(rec)) {
if (rec->opts.affinity != PERF_AFFINITY_SYS) {
pr_err("--affinity option is mutually exclusive to parallel streaming mode.\n");
goto out_opts;
}
if (record__aio_enabled(rec)) {
pr_err("Asynchronous streaming mode (--aio) is mutually exclusive to parallel streaming mode.\n");
goto out_opts;
}
}
if (rec->opts.comp_level != 0) {
pr_debug("Compression enabled, disabling build id collection at the end of the session.\n");
rec->no_buildid = true;
}
if (rec->opts.record_switch_events &&
!perf_can_record_switch_events()) {
ui__error("kernel does not support recording context switch events\n");
parse_options_usage(record_usage, record_options, "switch-events", 0);
err = -EINVAL;
goto out_opts;
}
if (rec->switch_output.time) {
signal(SIGALRM, alarm_sig_handler);
alarm(rec->switch_output.time);
}
if (rec->switch_output.num_files) {
rec->switch_output.filenames = calloc(rec->switch_output.num_files,
sizeof(char *));
if (!rec->switch_output.filenames) {
err = -EINVAL;
goto out_opts;
}
}
if (rec->timestamp_filename && record__threads_enabled(rec)) {
rec->timestamp_filename = false;
pr_warning("WARNING: --timestamp-filename option is not available in parallel streaming mode.\n");
}
/* For backward compatibility, -d implies --mem-info */
if (rec->opts.sample_address)
rec->opts.sample_data_src = true;
/*
* Allow aliases to facilitate the lookup of symbols for address
* filters. Refer to auxtrace_parse_filters().
*/
symbol_conf.allow_aliases = true;
symbol__init(NULL);
err = record__auxtrace_init(rec);
if (err)
goto out;
if (dry_run)
goto out;
err = -ENOMEM;
if (rec->no_buildid_cache || rec->no_buildid) {
disable_buildid_cache();
} else if (rec->switch_output.enabled) {
/*
* In 'perf record --switch-output', disable buildid
* generation by default to reduce data file switching
* overhead. Still generate buildid if they are required
* explicitly using
*
* perf record --switch-output --no-no-buildid \
* --no-no-buildid-cache
*
* Following code equals to:
*
* if ((rec->no_buildid || !rec->no_buildid_set) &&
* (rec->no_buildid_cache || !rec->no_buildid_cache_set))
* disable_buildid_cache();
*/
bool disable = true;
if (rec->no_buildid_set && !rec->no_buildid)
disable = false;
if (rec->no_buildid_cache_set && !rec->no_buildid_cache)
disable = false;
if (disable) {
rec->no_buildid = true;
rec->no_buildid_cache = true;
disable_buildid_cache();
}
}
if (record.opts.overwrite)
record.opts.tail_synthesize = true;
if (rec->evlist->core.nr_entries == 0) {
err = parse_event(rec->evlist, "cycles:P");
if (err)
goto out;
}
if (rec->opts.target.tid && !rec->opts.no_inherit_set)
rec->opts.no_inherit = true;
if (callchain_param.enabled && callchain_param.record_mode == CALLCHAIN_FP)
arch__add_leaf_frame_record_opts(&rec->opts);
err = -ENOMEM;
if (evlist__create_maps(rec->evlist, &rec->opts.target) < 0) {
if (rec->opts.target.pid != NULL) {
pr_err("Couldn't create thread/CPU maps: %s\n",
errno == ENOENT ? "No such process" : str_error_r(errno, errbuf, sizeof(errbuf)));
goto out;
}
else
usage_with_options(record_usage, record_options);
}
err = auxtrace_record__options(rec->itr, rec->evlist, &rec->opts);
if (err)
goto out;
/*
* We take all buildids when the file contains
* AUX area tracing data because we do not decode the
* trace because it would take too long.
*/
if (rec->opts.full_auxtrace)
rec->buildid_all = true;
if (rec->opts.text_poke) {
err = record__config_text_poke(rec->evlist);
if (err) {
pr_err("record__config_text_poke failed, error %d\n", err);
goto out;
}
}
if (rec->off_cpu) {
err = record__config_off_cpu(rec);
if (err) {
pr_err("record__config_off_cpu failed, error %d\n", err);
goto out;
}
}
if (record_opts__config(&rec->opts)) {
err = -EINVAL;
goto out;
}
Die Informationen auf dieser Webseite wurden
nach bestem Wissen sorgfältig zusammengestellt. Es wird jedoch weder Vollständigkeit, noch Richtigkeit,
noch Qualität der bereit gestellten Informationen zugesichert.
Bemerkung:
Die farbliche Syntaxdarstellung und die Messung sind noch experimentell.