1[
2    {
3        "BriefDescription": "Instructions Per Cycle (per Logical Processor)",
4        "MetricExpr": "INST_RETIRED.ANY / CPU_CLK_UNHALTED.THREAD",
5        "MetricGroup": "Summary",
6        "MetricName": "IPC"
7    },
8    {
9        "BriefDescription": "Uops Per Instruction",
10        "MetricExpr": "UOPS_RETIRED.RETIRE_SLOTS / INST_RETIRED.ANY",
11        "MetricGroup": "Pipeline;Retire",
12        "MetricName": "UPI"
13    },
14    {
15        "BriefDescription": "Instruction per taken branch",
16        "MetricExpr": "INST_RETIRED.ANY / BR_INST_RETIRED.NEAR_TAKEN",
17        "MetricGroup": "Branches;FetchBW;PGO",
18        "MetricName": "IpTB"
19    },
20    {
21        "BriefDescription": "Cycles Per Instruction (per Logical Processor)",
22        "MetricExpr": "1 / (INST_RETIRED.ANY / CPU_CLK_UNHALTED.THREAD)",
23        "MetricGroup": "Pipeline",
24        "MetricName": "CPI"
25    },
26    {
27        "BriefDescription": "Per-Logical Processor actual clocks when the Logical Processor is active.",
28        "MetricExpr": "CPU_CLK_UNHALTED.THREAD",
29        "MetricGroup": "Pipeline",
30        "MetricName": "CLKS"
31    },
32    {
33        "BriefDescription": "Instructions Per Cycle (per physical core)",
34        "MetricExpr": "INST_RETIRED.ANY / CPU_CLK_UNHALTED.THREAD",
35        "MetricGroup": "SMT;TmaL1",
36        "MetricName": "CoreIPC"
37    },
38    {
39        "BriefDescription": "Instructions Per Cycle (per physical core)",
40        "MetricExpr": "INST_RETIRED.ANY / ( ( CPU_CLK_UNHALTED.THREAD / 2 ) * ( 1 + CPU_CLK_UNHALTED.ONE_THREAD_ACTIVE / CPU_CLK_UNHALTED.REF_XCLK ) )",
41        "MetricGroup": "SMT;TmaL1",
42        "MetricName": "CoreIPC_SMT"
43    },
44    {
45        "BriefDescription": "Floating Point Operations Per Cycle",
46        "MetricExpr": "( 1 * ( FP_ARITH_INST_RETIRED.SCALAR_SINGLE + FP_ARITH_INST_RETIRED.SCALAR_DOUBLE ) + 2 * FP_ARITH_INST_RETIRED.128B_PACKED_DOUBLE + 4 * ( FP_ARITH_INST_RETIRED.128B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.256B_PACKED_DOUBLE ) + 8 * ( FP_ARITH_INST_RETIRED.256B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.512B_PACKED_DOUBLE ) + 16 * FP_ARITH_INST_RETIRED.512B_PACKED_SINGLE ) / CPU_CLK_UNHALTED.THREAD",
47        "MetricGroup": "Flops",
48        "MetricName": "FLOPc"
49    },
50    {
51        "BriefDescription": "Floating Point Operations Per Cycle",
52        "MetricExpr": "( 1 * ( FP_ARITH_INST_RETIRED.SCALAR_SINGLE + FP_ARITH_INST_RETIRED.SCALAR_DOUBLE ) + 2 * FP_ARITH_INST_RETIRED.128B_PACKED_DOUBLE + 4 * ( FP_ARITH_INST_RETIRED.128B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.256B_PACKED_DOUBLE ) + 8 * ( FP_ARITH_INST_RETIRED.256B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.512B_PACKED_DOUBLE ) + 16 * FP_ARITH_INST_RETIRED.512B_PACKED_SINGLE ) / ( ( CPU_CLK_UNHALTED.THREAD / 2 ) * ( 1 + CPU_CLK_UNHALTED.ONE_THREAD_ACTIVE / CPU_CLK_UNHALTED.REF_XCLK ) )",
53        "MetricGroup": "Flops_SMT",
54        "MetricName": "FLOPc_SMT"
55    },
56    {
57        "BriefDescription": "Instruction-Level-Parallelism (average number of uops executed when there is at least 1 uop executed)",
58        "MetricExpr": "UOPS_EXECUTED.THREAD / (( UOPS_EXECUTED.CORE_CYCLES_GE_1 / 2 ) if #SMT_on else UOPS_EXECUTED.CORE_CYCLES_GE_1)",
59        "MetricGroup": "Pipeline;PortsUtil",
60        "MetricName": "ILP"
61    },
62    {
63        "BriefDescription": "Number of Instructions per non-speculative Branch Misprediction (JEClear)",
64        "MetricExpr": "INST_RETIRED.ANY / BR_MISP_RETIRED.ALL_BRANCHES",
65        "MetricGroup": "BrMispredicts",
66        "MetricName": "IpMispredict"
67    },
68    {
69        "BriefDescription": "Core actual clocks when any Logical Processor is active on the Physical Core",
70        "MetricExpr": "( CPU_CLK_UNHALTED.THREAD_ANY / 2 ) if #SMT_on else CPU_CLK_UNHALTED.THREAD",
71        "MetricGroup": "SMT",
72        "MetricName": "CORE_CLKS"
73    },
74    {
75        "BriefDescription": "Instructions per Load (lower number means higher occurrence rate)",
76        "MetricExpr": "INST_RETIRED.ANY / MEM_INST_RETIRED.ALL_LOADS",
77        "MetricGroup": "InsType",
78        "MetricName": "IpLoad"
79    },
80    {
81        "BriefDescription": "Instructions per Store (lower number means higher occurrence rate)",
82        "MetricExpr": "INST_RETIRED.ANY / MEM_INST_RETIRED.ALL_STORES",
83        "MetricGroup": "InsType",
84        "MetricName": "IpStore"
85    },
86    {
87        "BriefDescription": "Instructions per Branch (lower number means higher occurrence rate)",
88        "MetricExpr": "INST_RETIRED.ANY / BR_INST_RETIRED.ALL_BRANCHES",
89        "MetricGroup": "Branches;InsType",
90        "MetricName": "IpBranch"
91    },
92    {
93        "BriefDescription": "Instructions per (near) call (lower number means higher occurrence rate)",
94        "MetricExpr": "INST_RETIRED.ANY / BR_INST_RETIRED.NEAR_CALL",
95        "MetricGroup": "Branches",
96        "MetricName": "IpCall"
97    },
98    {
99        "BriefDescription": "Branch instructions per taken branch. ",
100        "MetricExpr": "BR_INST_RETIRED.ALL_BRANCHES / BR_INST_RETIRED.NEAR_TAKEN",
101        "MetricGroup": "Branches;PGO",
102        "MetricName": "BpTkBranch"
103    },
104    {
105        "BriefDescription": "Instructions per Floating Point (FP) Operation (lower number means higher occurrence rate)",
106        "MetricExpr": "INST_RETIRED.ANY / ( 1 * ( FP_ARITH_INST_RETIRED.SCALAR_SINGLE + FP_ARITH_INST_RETIRED.SCALAR_DOUBLE ) + 2 * FP_ARITH_INST_RETIRED.128B_PACKED_DOUBLE + 4 * ( FP_ARITH_INST_RETIRED.128B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.256B_PACKED_DOUBLE ) + 8 * ( FP_ARITH_INST_RETIRED.256B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.512B_PACKED_DOUBLE ) + 16 * FP_ARITH_INST_RETIRED.512B_PACKED_SINGLE )",
107        "MetricGroup": "Flops;FpArith;InsType",
108        "MetricName": "IpFLOP"
109    },
110    {
111        "BriefDescription": "Total number of retired Instructions, Sample with: INST_RETIRED.PREC_DIST",
112        "MetricExpr": "INST_RETIRED.ANY",
113        "MetricGroup": "Summary;TmaL1",
114        "MetricName": "Instructions"
115    },
116    {
117        "BriefDescription": "Fraction of Uops delivered by the LSD (Loop Stream Detector; aka Loop Cache)",
118        "MetricExpr": "LSD.UOPS / (IDQ.DSB_UOPS + LSD.UOPS + IDQ.MITE_UOPS + IDQ.MS_UOPS)",
119        "MetricGroup": "LSD",
120        "MetricName": "LSD_Coverage"
121    },
122    {
123        "BriefDescription": "Fraction of Uops delivered by the DSB (aka Decoded ICache; or Uop Cache)",
124        "MetricExpr": "IDQ.DSB_UOPS / (IDQ.DSB_UOPS + LSD.UOPS + IDQ.MITE_UOPS + IDQ.MS_UOPS)",
125        "MetricGroup": "DSB;FetchBW",
126        "MetricName": "DSB_Coverage"
127    },
128    {
129        "BriefDescription": "Actual Average Latency for L1 data-cache miss demand loads (in core cycles)",
130        "MetricExpr": "L1D_PEND_MISS.PENDING / ( MEM_LOAD_RETIRED.L1_MISS + MEM_LOAD_RETIRED.FB_HIT )",
131        "MetricGroup": "MemoryBound;MemoryLat",
132        "MetricName": "Load_Miss_Real_Latency"
133    },
134    {
135        "BriefDescription": "Memory-Level-Parallelism (average number of L1 miss demand load when there is at least one such miss. Per-Logical Processor)",
136        "MetricExpr": "L1D_PEND_MISS.PENDING / L1D_PEND_MISS.PENDING_CYCLES",
137        "MetricGroup": "MemoryBound;MemoryBW",
138        "MetricName": "MLP"
139    },
140    {
141        "BriefDescription": "Utilization of the core's Page Walker(s) serving STLB misses triggered by instruction/Load/Store accesses",
142        "MetricConstraint": "NO_NMI_WATCHDOG",
143        "MetricExpr": "( ITLB_MISSES.WALK_PENDING + DTLB_LOAD_MISSES.WALK_PENDING + DTLB_STORE_MISSES.WALK_PENDING + EPT.WALK_PENDING ) / ( 2 * CORE_CLKS )",
144        "MetricGroup": "MemoryTLB",
145        "MetricName": "Page_Walks_Utilization"
146    },
147    {
148        "BriefDescription": "Average data fill bandwidth to the L1 data cache [GB / sec]",
149        "MetricExpr": "64 * L1D.REPLACEMENT / 1000000000 / duration_time",
150        "MetricGroup": "MemoryBW",
151        "MetricName": "L1D_Cache_Fill_BW"
152    },
153    {
154        "BriefDescription": "Average data fill bandwidth to the L2 cache [GB / sec]",
155        "MetricExpr": "64 * L2_LINES_IN.ALL / 1000000000 / duration_time",
156        "MetricGroup": "MemoryBW",
157        "MetricName": "L2_Cache_Fill_BW"
158    },
159    {
160        "BriefDescription": "Average per-core data fill bandwidth to the L3 cache [GB / sec]",
161        "MetricExpr": "64 * LONGEST_LAT_CACHE.MISS / 1000000000 / duration_time",
162        "MetricGroup": "MemoryBW",
163        "MetricName": "L3_Cache_Fill_BW"
164    },
165    {
166        "BriefDescription": "Average per-core data access bandwidth to the L3 cache [GB / sec]",
167        "MetricExpr": "64 * OFFCORE_REQUESTS.ALL_REQUESTS / 1000000000 / duration_time",
168        "MetricGroup": "MemoryBW;Offcore",
169        "MetricName": "L3_Cache_Access_BW"
170    },
171    {
172        "BriefDescription": "L1 cache true misses per kilo instruction for retired demand loads",
173        "MetricExpr": "1000 * MEM_LOAD_RETIRED.L1_MISS / INST_RETIRED.ANY",
174        "MetricGroup": "CacheMisses",
175        "MetricName": "L1MPKI"
176    },
177    {
178        "BriefDescription": "L2 cache true misses per kilo instruction for retired demand loads",
179        "MetricExpr": "1000 * MEM_LOAD_RETIRED.L2_MISS / INST_RETIRED.ANY",
180        "MetricGroup": "CacheMisses",
181        "MetricName": "L2MPKI"
182    },
183    {
184        "BriefDescription": "L2 cache misses per kilo instruction for all request types (including speculative)",
185        "MetricExpr": "1000 * L2_RQSTS.MISS / INST_RETIRED.ANY",
186        "MetricGroup": "CacheMisses;Offcore",
187        "MetricName": "L2MPKI_All"
188    },
189    {
190        "BriefDescription": "L2 cache hits per kilo instruction for all request types (including speculative)",
191        "MetricExpr": "1000 * ( L2_RQSTS.REFERENCES - L2_RQSTS.MISS ) / INST_RETIRED.ANY",
192        "MetricGroup": "CacheMisses",
193        "MetricName": "L2HPKI_All"
194    },
195    {
196        "BriefDescription": "L3 cache true misses per kilo instruction for retired demand loads",
197        "MetricExpr": "1000 * MEM_LOAD_RETIRED.L3_MISS / INST_RETIRED.ANY",
198        "MetricGroup": "CacheMisses",
199        "MetricName": "L3MPKI"
200    },
201    {
202        "BriefDescription": "Rate of silent evictions from the L2 cache per Kilo instruction where the evicted lines are dropped (no writeback to L3 or memory)",
203        "MetricExpr": "1000 * L2_LINES_OUT.SILENT / INST_RETIRED.ANY",
204        "MetricGroup": "L2Evicts;Server",
205        "MetricName": "L2_Evictions_Silent_PKI"
206    },
207    {
208        "BriefDescription": "Rate of non silent evictions from the L2 cache per Kilo instruction",
209        "MetricExpr": "1000 * L2_LINES_OUT.NON_SILENT / INST_RETIRED.ANY",
210        "MetricGroup": "L2Evicts;Server",
211        "MetricName": "L2_Evictions_NonSilent_PKI"
212    },
213    {
214        "BriefDescription": "Average CPU Utilization",
215        "MetricExpr": "CPU_CLK_UNHALTED.REF_TSC / msr@tsc@",
216        "MetricGroup": "HPC;Summary",
217        "MetricName": "CPU_Utilization"
218    },
219    {
220        "BriefDescription": "Measured Average Frequency for unhalted processors [GHz]",
221        "MetricExpr": "(CPU_CLK_UNHALTED.THREAD / CPU_CLK_UNHALTED.REF_TSC) * msr@tsc@ / 1000000000 / duration_time",
222        "MetricGroup": "Summary;Power",
223        "MetricName": "Average_Frequency"
224    },
225    {
226        "BriefDescription": "Giga Floating Point Operations Per Second",
227        "MetricExpr": "( ( 1 * ( FP_ARITH_INST_RETIRED.SCALAR_SINGLE + FP_ARITH_INST_RETIRED.SCALAR_DOUBLE ) + 2 * FP_ARITH_INST_RETIRED.128B_PACKED_DOUBLE + 4 * ( FP_ARITH_INST_RETIRED.128B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.256B_PACKED_DOUBLE ) + 8 * ( FP_ARITH_INST_RETIRED.256B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.512B_PACKED_DOUBLE ) + 16 * FP_ARITH_INST_RETIRED.512B_PACKED_SINGLE ) / 1000000000 ) / duration_time",
228        "MetricGroup": "Flops;HPC",
229        "MetricName": "GFLOPs"
230    },
231    {
232        "BriefDescription": "Average Frequency Utilization relative nominal frequency",
233        "MetricExpr": "CPU_CLK_UNHALTED.THREAD / CPU_CLK_UNHALTED.REF_TSC",
234        "MetricGroup": "Power",
235        "MetricName": "Turbo_Utilization"
236    },
237    {
238        "BriefDescription": "Fraction of cycles where both hardware Logical Processors were active",
239        "MetricExpr": "1 - CPU_CLK_UNHALTED.ONE_THREAD_ACTIVE / ( CPU_CLK_UNHALTED.REF_XCLK_ANY / 2 ) if #SMT_on else 0",
240        "MetricGroup": "SMT",
241        "MetricName": "SMT_2T_Utilization"
242    },
243    {
244        "BriefDescription": "Fraction of cycles spent in the Operating System (OS) Kernel mode",
245        "MetricExpr": "CPU_CLK_UNHALTED.THREAD_P:k / CPU_CLK_UNHALTED.THREAD",
246        "MetricGroup": "OS",
247        "MetricName": "Kernel_Utilization"
248    },
249    {
250        "BriefDescription": "Average external Memory Bandwidth Use for reads and writes [GB / sec]",
251        "MetricExpr": "( 64 * ( uncore_imc@cas_count_read@ + uncore_imc@cas_count_write@ ) / 1000000000 ) / duration_time",
252        "MetricGroup": "HPC;MemoryBW;SoC",
253        "MetricName": "DRAM_BW_Use"
254    },
255    {
256        "BriefDescription": "Average latency of data read request to external memory (in nanoseconds). Accounts for demand loads and L1/L2 prefetches",
257        "MetricExpr": "1000000000 * ( cha@event\\=0x36\\,umask\\=0x21\\,config\\=0x40433@ / cha@event\\=0x35\\,umask\\=0x21\\,config\\=0x40433@ ) / ( cha_0@event\\=0x0@ / duration_time )",
258        "MetricGroup": "MemoryLat;SoC",
259        "MetricName": "MEM_Read_Latency"
260    },
261    {
262        "BriefDescription": "Average number of parallel data read requests to external memory. Accounts for demand loads and L1/L2 prefetches",
263        "MetricExpr": "cha@event\\=0x36\\,umask\\=0x21\\,config\\=0x40433@ / cha@event\\=0x36\\,umask\\=0x21\\,config\\=0x40433\\,thresh\\=1@",
264        "MetricGroup": "MemoryBW;SoC",
265        "MetricName": "MEM_Parallel_Reads"
266    },
267    {
268        "BriefDescription": "Average latency of data read request to external 3D X-Point memory [in nanoseconds]. Accounts for demand loads and L1/L2 data-read prefetches",
269        "MetricExpr": "( 1000000000 * ( imc@event\\=0xe0\\,umask\\=0x1@ / imc@event\\=0xe3@ ) / imc_0@event\\=0x0@ )",
270        "MetricGroup": "MemoryLat;SoC;Server",
271        "MetricName": "MEM_PMM_Read_Latency"
272    },
273    {
274        "BriefDescription": "Average 3DXP Memory Bandwidth Use for reads [GB / sec]",
275        "MetricExpr": "( ( 64 * imc@event\\=0xe3@ / 1000000000 ) / duration_time )",
276        "MetricGroup": "MemoryBW;SoC;Server",
277        "MetricName": "PMM_Read_BW"
278    },
279    {
280        "BriefDescription": "Average 3DXP Memory Bandwidth Use for Writes [GB / sec]",
281        "MetricExpr": "( ( 64 * imc@event\\=0xe7@ / 1000000000 ) / duration_time )",
282        "MetricGroup": "MemoryBW;SoC;Server",
283        "MetricName": "PMM_Write_BW"
284    },
285    {
286        "BriefDescription": "Average IO (network or disk) Bandwidth Use for Writes [GB / sec]",
287        "MetricExpr": "( UNC_IIO_DATA_REQ_OF_CPU.MEM_READ.PART0 + UNC_IIO_DATA_REQ_OF_CPU.MEM_READ.PART1 + UNC_IIO_DATA_REQ_OF_CPU.MEM_READ.PART2 + UNC_IIO_DATA_REQ_OF_CPU.MEM_READ.PART3 ) * 4 / 1000000000 / duration_time",
288        "MetricGroup": "IoBW;SoC;Server",
289        "MetricName": "IO_Write_BW"
290    },
291    {
292        "BriefDescription": "Average IO (network or disk) Bandwidth Use for Reads [GB / sec]",
293        "MetricExpr": "( UNC_IIO_DATA_REQ_OF_CPU.MEM_WRITE.PART0 + UNC_IIO_DATA_REQ_OF_CPU.MEM_WRITE.PART1 + UNC_IIO_DATA_REQ_OF_CPU.MEM_WRITE.PART2 + UNC_IIO_DATA_REQ_OF_CPU.MEM_WRITE.PART3 ) * 4 / 1000000000 / duration_time",
294        "MetricGroup": "IoBW;SoC;Server",
295        "MetricName": "IO_Read_BW"
296    },
297    {
298        "BriefDescription": "Socket actual clocks when any core is active on that socket",
299        "MetricExpr": "cha_0@event\\=0x0@",
300        "MetricGroup": "SoC",
301        "MetricName": "Socket_CLKS"
302    },
303    {
304        "BriefDescription": "Instructions per Far Branch ( Far Branches apply upon transition from application to operating system, handling interrupts, exceptions) [lower number means higher occurrence rate]",
305        "MetricExpr": "INST_RETIRED.ANY / BR_INST_RETIRED.FAR_BRANCH:u",
306        "MetricGroup": "Branches;OS",
307        "MetricName": "IpFarBranch"
308    },
309    {
310        "BriefDescription": "C3 residency percent per core",
311        "MetricExpr": "(cstate_core@c3\\-residency@ / msr@tsc@) * 100",
312        "MetricGroup": "Power",
313        "MetricName": "C3_Core_Residency"
314    },
315    {
316        "BriefDescription": "C6 residency percent per core",
317        "MetricExpr": "(cstate_core@c6\\-residency@ / msr@tsc@) * 100",
318        "MetricGroup": "Power",
319        "MetricName": "C6_Core_Residency"
320    },
321    {
322        "BriefDescription": "C7 residency percent per core",
323        "MetricExpr": "(cstate_core@c7\\-residency@ / msr@tsc@) * 100",
324        "MetricGroup": "Power",
325        "MetricName": "C7_Core_Residency"
326    },
327    {
328        "BriefDescription": "C2 residency percent per package",
329        "MetricExpr": "(cstate_pkg@c2\\-residency@ / msr@tsc@) * 100",
330        "MetricGroup": "Power",
331        "MetricName": "C2_Pkg_Residency"
332    },
333    {
334        "BriefDescription": "C3 residency percent per package",
335        "MetricExpr": "(cstate_pkg@c3\\-residency@ / msr@tsc@) * 100",
336        "MetricGroup": "Power",
337        "MetricName": "C3_Pkg_Residency"
338    },
339    {
340        "BriefDescription": "C6 residency percent per package",
341        "MetricExpr": "(cstate_pkg@c6\\-residency@ / msr@tsc@) * 100",
342        "MetricGroup": "Power",
343        "MetricName": "C6_Pkg_Residency"
344    },
345    {
346        "BriefDescription": "C7 residency percent per package",
347        "MetricExpr": "(cstate_pkg@c7\\-residency@ / msr@tsc@) * 100",
348        "MetricGroup": "Power",
349        "MetricName": "C7_Pkg_Residency"
350    }
351]
352