aboutsummaryrefslogtreecommitdiffstats
path: root/arch/powerpc/kernel
diff options
context:
space:
mode:
authorPaul Mackerras <paulus@samba.org>2009-01-14 05:00:30 -0500
committerPaul Mackerras <paulus@samba.org>2009-01-14 05:00:30 -0500
commit3b6f9e5cb21964b7ce12bf81076f830885563ec8 (patch)
treee9d5ecffafa66cc3aeb259ade15a2611ad795327 /arch/powerpc/kernel
parent01d0287f068de2934109ba9b989d8807526cccc2 (diff)
perf_counter: Add support for pinned and exclusive counter groups
Impact: New perf_counter features A pinned counter group is one that the user wants to have on the CPU whenever possible, i.e. whenever the associated task is running, for a per-task group, or always for a per-cpu group. If the system cannot satisfy that, it puts the group into an error state where it is not scheduled any more and reads from it return EOF (i.e. 0 bytes read). The group can be released from error state and made readable again using prctl(PR_TASK_PERF_COUNTERS_ENABLE). When we have finer-grained enable/disable controls on counters we'll be able to reset the error state on individual groups. An exclusive group is one that the user wants to be the only group using the CPU performance monitor hardware whenever it is on. The counter group scheduler will not schedule an exclusive group if there are already other groups on the CPU and will not schedule other groups onto the CPU if there is an exclusive group scheduled (that statement does not apply to groups containing only software counters, which can always go on and which do not prevent an exclusive group from going on). With an exclusive group, we will be able to let users program PMU registers at a low level without the concern that those settings will perturb other measurements. Along the way this reorganizes things a little: - is_software_counter() is moved to perf_counter.h. - cpuctx->active_oncpu now records the number of hardware counters on the CPU, i.e. it now excludes software counters. Nothing was reading cpuctx->active_oncpu before, so this change is harmless. - A new cpuctx->exclusive field records whether we currently have an exclusive group on the CPU. - counter_sched_out moves higher up in perf_counter.c and gets called from __perf_counter_remove_from_context and __perf_counter_exit_task, where we used to have essentially the same code. - __perf_counter_sched_in now goes through the counter list twice, doing the pinned counters in the first loop and the non-pinned counters in the second loop, in order to give the pinned counters the best chance to be scheduled in. Note that only a group leader can be exclusive or pinned, and that attribute applies to the whole group. This avoids some awkwardness in some corner cases (e.g. where a group leader is closed and the other group members get added to the context list). If we want to relax that restriction later, we can, and it is easier to relax a restriction than to apply a new one. This doesn't yet handle the case where a pinned counter is inherited and goes into error state in the child - the error state is not propagated up to the parent when the child exits, and arguably it should. Signed-off-by: Paul Mackerras <paulus@samba.org>
Diffstat (limited to 'arch/powerpc/kernel')
-rw-r--r--arch/powerpc/kernel/perf_counter.c10
1 files changed, 1 insertions, 9 deletions
diff --git a/arch/powerpc/kernel/perf_counter.c b/arch/powerpc/kernel/perf_counter.c
index 85ad25923c2c..5b0211348c73 100644
--- a/arch/powerpc/kernel/perf_counter.c
+++ b/arch/powerpc/kernel/perf_counter.c
@@ -36,14 +36,6 @@ void perf_counter_print_debug(void)
36} 36}
37 37
38/* 38/*
39 * Return 1 for a software counter, 0 for a hardware counter
40 */
41static inline int is_software_counter(struct perf_counter *counter)
42{
43 return !counter->hw_event.raw && counter->hw_event.type < 0;
44}
45
46/*
47 * Read one performance monitor counter (PMC). 39 * Read one performance monitor counter (PMC).
48 */ 40 */
49static unsigned long read_pmc(int idx) 41static unsigned long read_pmc(int idx)
@@ -443,6 +435,7 @@ int hw_perf_group_sched_in(struct perf_counter *group_leader,
443 */ 435 */
444 for (i = n0; i < n0 + n; ++i) 436 for (i = n0; i < n0 + n; ++i)
445 cpuhw->counter[i]->hw.config = cpuhw->events[i]; 437 cpuhw->counter[i]->hw.config = cpuhw->events[i];
438 cpuctx->active_oncpu += n;
446 n = 1; 439 n = 1;
447 counter_sched_in(group_leader, cpu); 440 counter_sched_in(group_leader, cpu);
448 list_for_each_entry(sub, &group_leader->sibling_list, list_entry) { 441 list_for_each_entry(sub, &group_leader->sibling_list, list_entry) {
@@ -451,7 +444,6 @@ int hw_perf_group_sched_in(struct perf_counter *group_leader,
451 ++n; 444 ++n;
452 } 445 }
453 } 446 }
454 cpuctx->active_oncpu += n;
455 ctx->nr_active += n; 447 ctx->nr_active += n;
456 448
457 return 1; 449 return 1;