3 "BriefDescription": "Instructions Per Cycle (per Logical Processor)",
4 "MetricExpr": "INST_RETIRED.ANY / CPU_CLK_UNHALTED.THREAD",
5 "MetricGroup": "Summary",
9 "BriefDescription": "Uops Per Instruction",
10 "MetricExpr": "UOPS_RETIRED.SLOTS / INST_RETIRED.ANY",
11 "MetricGroup": "Pipeline;Retire",
15 "BriefDescription": "Instruction per taken branch",
16 "MetricExpr": "INST_RETIRED.ANY / BR_INST_RETIRED.NEAR_TAKEN",
17 "MetricGroup": "Branches;FetchBW;PGO",
21 "BriefDescription": "Cycles Per Instruction (per Logical Processor)",
22 "MetricExpr": "1 / (INST_RETIRED.ANY / CPU_CLK_UNHALTED.THREAD)",
23 "MetricGroup": "Pipeline",
27 "BriefDescription": "Per-Logical Processor actual clocks when the Logical Processor is active.",
28 "MetricExpr": "CPU_CLK_UNHALTED.THREAD",
29 "MetricGroup": "Pipeline",
33 "BriefDescription": "Instructions Per Cycle (per physical core)",
34 "MetricExpr": "INST_RETIRED.ANY / CPU_CLK_UNHALTED.DISTRIBUTED",
35 "MetricGroup": "SMT;TmaL1",
36 "MetricName": "CoreIPC"
39 "BriefDescription": "Floating Point Operations Per Cycle",
40 "MetricExpr": "( 1 * ( FP_ARITH_INST_RETIRED.SCALAR_SINGLE + FP_ARITH_INST_RETIRED.SCALAR_DOUBLE ) + 2 * FP_ARITH_INST_RETIRED.128B_PACKED_DOUBLE + 4 * ( FP_ARITH_INST_RETIRED.128B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.256B_PACKED_DOUBLE ) + 8 * ( FP_ARITH_INST_RETIRED.256B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.512B_PACKED_DOUBLE ) + 16 * FP_ARITH_INST_RETIRED.512B_PACKED_SINGLE ) / CPU_CLK_UNHALTED.DISTRIBUTED",
41 "MetricGroup": "Flops",
45 "BriefDescription": "Instruction-Level-Parallelism (average number of uops executed when there is at least 1 uop executed)",
46 "MetricExpr": "UOPS_EXECUTED.THREAD / (( UOPS_EXECUTED.CORE_CYCLES_GE_1 / 2 ) if #SMT_on else UOPS_EXECUTED.CORE_CYCLES_GE_1)",
47 "MetricGroup": "Pipeline;PortsUtil",
51 "BriefDescription": "Number of Instructions per non-speculative Branch Misprediction (JEClear)",
52 "MetricExpr": "INST_RETIRED.ANY / BR_MISP_RETIRED.ALL_BRANCHES",
53 "MetricGroup": "BrMispredicts",
54 "MetricName": "IpMispredict"
57 "BriefDescription": "Core actual clocks when any Logical Processor is active on the Physical Core",
58 "MetricExpr": "CPU_CLK_UNHALTED.DISTRIBUTED",
60 "MetricName": "CORE_CLKS"
63 "BriefDescription": "Instructions per Load (lower number means higher occurrence rate)",
64 "MetricExpr": "INST_RETIRED.ANY / MEM_INST_RETIRED.ALL_LOADS",
65 "MetricGroup": "InsType",
66 "MetricName": "IpLoad"
69 "BriefDescription": "Instructions per Store (lower number means higher occurrence rate)",
70 "MetricExpr": "INST_RETIRED.ANY / MEM_INST_RETIRED.ALL_STORES",
71 "MetricGroup": "InsType",
72 "MetricName": "IpStore"
75 "BriefDescription": "Instructions per Branch (lower number means higher occurrence rate)",
76 "MetricExpr": "INST_RETIRED.ANY / BR_INST_RETIRED.ALL_BRANCHES",
77 "MetricGroup": "Branches;InsType",
78 "MetricName": "IpBranch"
81 "BriefDescription": "Instructions per (near) call (lower number means higher occurrence rate)",
82 "MetricExpr": "INST_RETIRED.ANY / BR_INST_RETIRED.NEAR_CALL",
83 "MetricGroup": "Branches",
84 "MetricName": "IpCall"
87 "BriefDescription": "Branch instructions per taken branch. ",
88 "MetricExpr": "BR_INST_RETIRED.ALL_BRANCHES / BR_INST_RETIRED.NEAR_TAKEN",
89 "MetricGroup": "Branches;PGO",
90 "MetricName": "BpTkBranch"
93 "BriefDescription": "Instructions per Floating Point (FP) Operation (lower number means higher occurrence rate)",
94 "MetricExpr": "INST_RETIRED.ANY / ( 1 * ( FP_ARITH_INST_RETIRED.SCALAR_SINGLE + FP_ARITH_INST_RETIRED.SCALAR_DOUBLE ) + 2 * FP_ARITH_INST_RETIRED.128B_PACKED_DOUBLE + 4 * ( FP_ARITH_INST_RETIRED.128B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.256B_PACKED_DOUBLE ) + 8 * ( FP_ARITH_INST_RETIRED.256B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.512B_PACKED_DOUBLE ) + 16 * FP_ARITH_INST_RETIRED.512B_PACKED_SINGLE )",
95 "MetricGroup": "Flops;FpArith;InsType",
96 "MetricName": "IpFLOP"
99 "BriefDescription": "Total number of retired Instructions, Sample with: INST_RETIRED.PREC_DIST",
100 "MetricExpr": "INST_RETIRED.ANY",
101 "MetricGroup": "Summary;TmaL1",
102 "MetricName": "Instructions"
105 "BriefDescription": "Fraction of Uops delivered by the LSD (Loop Stream Detector; aka Loop Cache)",
106 "MetricExpr": "LSD.UOPS / (IDQ.DSB_UOPS + LSD.UOPS + IDQ.MITE_UOPS + IDQ.MS_UOPS)",
107 "MetricGroup": "LSD",
108 "MetricName": "LSD_Coverage"
111 "BriefDescription": "Fraction of Uops delivered by the DSB (aka Decoded ICache; or Uop Cache)",
112 "MetricExpr": "IDQ.DSB_UOPS / (IDQ.DSB_UOPS + LSD.UOPS + IDQ.MITE_UOPS + IDQ.MS_UOPS)",
113 "MetricGroup": "DSB;FetchBW",
114 "MetricName": "DSB_Coverage"
117 "BriefDescription": "Actual Average Latency for L1 data-cache miss demand loads (in core cycles)",
118 "MetricExpr": "L1D_PEND_MISS.PENDING / ( MEM_LOAD_RETIRED.L1_MISS + MEM_LOAD_RETIRED.FB_HIT )",
119 "MetricGroup": "MemoryBound;MemoryLat",
120 "MetricName": "Load_Miss_Real_Latency"
123 "BriefDescription": "Memory-Level-Parallelism (average number of L1 miss demand load when there is at least one such miss. Per-Logical Processor)",
124 "MetricExpr": "L1D_PEND_MISS.PENDING / L1D_PEND_MISS.PENDING_CYCLES",
125 "MetricGroup": "MemoryBound;MemoryBW",
129 "BriefDescription": "Utilization of the core's Page Walker(s) serving STLB misses triggered by instruction/Load/Store accesses",
130 "MetricConstraint": "NO_NMI_WATCHDOG",
131 "MetricExpr": "( ITLB_MISSES.WALK_PENDING + DTLB_LOAD_MISSES.WALK_PENDING + DTLB_STORE_MISSES.WALK_PENDING ) / ( 2 * CPU_CLK_UNHALTED.DISTRIBUTED )",
132 "MetricGroup": "MemoryTLB",
133 "MetricName": "Page_Walks_Utilization"
136 "BriefDescription": "Average data fill bandwidth to the L1 data cache [GB / sec]",
137 "MetricExpr": "64 * L1D.REPLACEMENT / 1000000000 / duration_time",
138 "MetricGroup": "MemoryBW",
139 "MetricName": "L1D_Cache_Fill_BW"
142 "BriefDescription": "Average data fill bandwidth to the L2 cache [GB / sec]",
143 "MetricExpr": "64 * L2_LINES_IN.ALL / 1000000000 / duration_time",
144 "MetricGroup": "MemoryBW",
145 "MetricName": "L2_Cache_Fill_BW"
148 "BriefDescription": "Average per-core data fill bandwidth to the L3 cache [GB / sec]",
149 "MetricExpr": "64 * LONGEST_LAT_CACHE.MISS / 1000000000 / duration_time",
150 "MetricGroup": "MemoryBW",
151 "MetricName": "L3_Cache_Fill_BW"
154 "BriefDescription": "Average per-core data access bandwidth to the L3 cache [GB / sec]",
155 "MetricExpr": "64 * OFFCORE_REQUESTS.ALL_REQUESTS / 1000000000 / duration_time",
156 "MetricGroup": "MemoryBW;Offcore",
157 "MetricName": "L3_Cache_Access_BW"
160 "BriefDescription": "L1 cache true misses per kilo instruction for retired demand loads",
161 "MetricExpr": "1000 * MEM_LOAD_RETIRED.L1_MISS / INST_RETIRED.ANY",
162 "MetricGroup": "CacheMisses",
163 "MetricName": "L1MPKI"
166 "BriefDescription": "L2 cache true misses per kilo instruction for retired demand loads",
167 "MetricExpr": "1000 * MEM_LOAD_RETIRED.L2_MISS / INST_RETIRED.ANY",
168 "MetricGroup": "CacheMisses",
169 "MetricName": "L2MPKI"
172 "BriefDescription": "L2 cache misses per kilo instruction for all request types (including speculative)",
173 "MetricExpr": "1000 * ( ( OFFCORE_REQUESTS.ALL_DATA_RD - OFFCORE_REQUESTS.DEMAND_DATA_RD ) + L2_RQSTS.ALL_DEMAND_MISS + L2_RQSTS.SWPF_MISS ) / INST_RETIRED.ANY",
174 "MetricGroup": "CacheMisses;Offcore",
175 "MetricName": "L2MPKI_All"
178 "BriefDescription": "L3 cache true misses per kilo instruction for retired demand loads",
179 "MetricExpr": "1000 * MEM_LOAD_RETIRED.L3_MISS / INST_RETIRED.ANY",
180 "MetricGroup": "CacheMisses",
181 "MetricName": "L3MPKI"
184 "BriefDescription": "Rate of silent evictions from the L2 cache per Kilo instruction where the evicted lines are dropped (no writeback to L3 or memory)",
185 "MetricExpr": "1000 * L2_LINES_OUT.SILENT / INST_RETIRED.ANY",
186 "MetricGroup": "L2Evicts;Server",
187 "MetricName": "L2_Evictions_Silent_PKI"
190 "BriefDescription": "Rate of non silent evictions from the L2 cache per Kilo instruction",
191 "MetricExpr": "1000 * L2_LINES_OUT.NON_SILENT / INST_RETIRED.ANY",
192 "MetricGroup": "L2Evicts;Server",
193 "MetricName": "L2_Evictions_NonSilent_PKI"
196 "BriefDescription": "Average CPU Utilization",
197 "MetricExpr": "CPU_CLK_UNHALTED.REF_TSC / msr@tsc@",
198 "MetricGroup": "HPC;Summary",
199 "MetricName": "CPU_Utilization"
202 "BriefDescription": "Measured Average Frequency for unhalted processors [GHz]",
203 "MetricExpr": "(CPU_CLK_UNHALTED.THREAD / CPU_CLK_UNHALTED.REF_TSC) * msr@tsc@ / 1000000000 / duration_time",
204 "MetricGroup": "Summary;Power",
205 "MetricName": "Average_Frequency"
208 "BriefDescription": "Giga Floating Point Operations Per Second",
209 "MetricExpr": "( ( 1 * ( FP_ARITH_INST_RETIRED.SCALAR_SINGLE + FP_ARITH_INST_RETIRED.SCALAR_DOUBLE ) + 2 * FP_ARITH_INST_RETIRED.128B_PACKED_DOUBLE + 4 * ( FP_ARITH_INST_RETIRED.128B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.256B_PACKED_DOUBLE ) + 8 * ( FP_ARITH_INST_RETIRED.256B_PACKED_SINGLE + FP_ARITH_INST_RETIRED.512B_PACKED_DOUBLE ) + 16 * FP_ARITH_INST_RETIRED.512B_PACKED_SINGLE ) / 1000000000 ) / duration_time",
210 "MetricGroup": "Flops;HPC",
211 "MetricName": "GFLOPs"
214 "BriefDescription": "Average Frequency Utilization relative nominal frequency",
215 "MetricExpr": "CPU_CLK_UNHALTED.THREAD / CPU_CLK_UNHALTED.REF_TSC",
216 "MetricGroup": "Power",
217 "MetricName": "Turbo_Utilization"
220 "BriefDescription": "Fraction of cycles where both hardware Logical Processors were active",
221 "MetricExpr": "1 - CPU_CLK_UNHALTED.ONE_THREAD_ACTIVE / CPU_CLK_UNHALTED.REF_DISTRIBUTED if #SMT_on else 0",
222 "MetricGroup": "SMT",
223 "MetricName": "SMT_2T_Utilization"
226 "BriefDescription": "Fraction of cycles spent in the Operating System (OS) Kernel mode",
227 "MetricExpr": "CPU_CLK_UNHALTED.THREAD_P:k / CPU_CLK_UNHALTED.THREAD",
229 "MetricName": "Kernel_Utilization"
232 "BriefDescription": "Average external Memory Bandwidth Use for reads and writes [GB / sec]",
233 "MetricExpr": "( 64 * ( uncore_imc@cas_count_read@ + uncore_imc@cas_count_write@ ) / 1000000000 ) / duration_time",
234 "MetricGroup": "HPC;MemoryBW;SoC",
235 "MetricName": "DRAM_BW_Use"
238 "BriefDescription": "Average latency of data read request to external memory (in nanoseconds). Accounts for demand loads and L1/L2 prefetches",
239 "MetricExpr": "1000000000 * ( UNC_CHA_TOR_OCCUPANCY.IA_MISS_DRD / UNC_CHA_TOR_INSERTS.IA_MISS_DRD ) / ( cha_0@event\\=0x0@ / duration_time )",
240 "MetricGroup": "MemoryLat;SoC",
241 "MetricName": "MEM_Read_Latency"
244 "BriefDescription": "Average number of parallel data read requests to external memory. Accounts for demand loads and L1/L2 prefetches",
245 "MetricExpr": "UNC_CHA_TOR_OCCUPANCY.IA_MISS_DRD / cha@event\\=0x36\\,umask\\=0xC817FE01\\,thresh\\=1@",
246 "MetricGroup": "MemoryBW;SoC",
247 "MetricName": "MEM_Parallel_Reads"
250 "BriefDescription": "Average latency of data read request to external 3D X-Point memory [in nanoseconds]. Accounts for demand loads and L1/L2 data-read prefetches",
251 "MetricExpr": "( 1000000000 * ( UNC_CHA_TOR_OCCUPANCY.IA_MISS_DRD_PMM / UNC_CHA_TOR_INSERTS.IA_MISS_DRD_PMM ) / cha_0@event\\=0x0@ )",
252 "MetricGroup": "MemoryLat;SoC;Server",
253 "MetricName": "MEM_PMM_Read_Latency"
256 "BriefDescription": "Average 3DXP Memory Bandwidth Use for reads [GB / sec]",
257 "MetricExpr": "( ( 64 * imc@event\\=0xe3@ / 1000000000 ) / duration_time )",
258 "MetricGroup": "MemoryBW;SoC;Server",
259 "MetricName": "PMM_Read_BW"
262 "BriefDescription": "Average 3DXP Memory Bandwidth Use for Writes [GB / sec]",
263 "MetricExpr": "( ( 64 * imc@event\\=0xe7@ / 1000000000 ) / duration_time )",
264 "MetricGroup": "MemoryBW;SoC;Server",
265 "MetricName": "PMM_Write_BW"
268 "BriefDescription": "Average IO (network or disk) Bandwidth Use for Writes [GB / sec]",
269 "MetricExpr": "UNC_CHA_TOR_INSERTS.IO_PCIRDCUR * 64 / 1000000000 / duration_time",
270 "MetricGroup": "IoBW;SoC;Server",
271 "MetricName": "IO_Write_BW"
274 "BriefDescription": "Average IO (network or disk) Bandwidth Use for Reads [GB / sec]",
275 "MetricExpr": "( UNC_CHA_TOR_INSERTS.IO_HIT_ITOM + UNC_CHA_TOR_INSERTS.IO_MISS_ITOM + UNC_CHA_TOR_INSERTS.IO_HIT_ITOMCACHENEAR + UNC_CHA_TOR_INSERTS.IO_MISS_ITOMCACHENEAR ) * 64 / 1000000000 / duration_time",
276 "MetricGroup": "IoBW;SoC;Server",
277 "MetricName": "IO_Read_BW"
280 "BriefDescription": "Socket actual clocks when any core is active on that socket",
281 "MetricExpr": "cha_0@event\\=0x0@",
282 "MetricGroup": "SoC",
283 "MetricName": "Socket_CLKS"
286 "BriefDescription": "Instructions per Far Branch ( Far Branches apply upon transition from application to operating system, handling interrupts, exceptions) [lower number means higher occurrence rate]",
287 "MetricExpr": "INST_RETIRED.ANY / BR_INST_RETIRED.FAR_BRANCH:u",
288 "MetricGroup": "Branches;OS",
289 "MetricName": "IpFarBranch"
292 "BriefDescription": "C1 residency percent per core",
293 "MetricExpr": "(cstate_core@c1\\-residency@ / msr@tsc@) * 100",
294 "MetricGroup": "Power",
295 "MetricName": "C1_Core_Residency"
298 "BriefDescription": "C6 residency percent per core",
299 "MetricExpr": "(cstate_core@c6\\-residency@ / msr@tsc@) * 100",
300 "MetricGroup": "Power",
301 "MetricName": "C6_Core_Residency"
304 "BriefDescription": "C2 residency percent per package",
305 "MetricExpr": "(cstate_pkg@c2\\-residency@ / msr@tsc@) * 100",
306 "MetricGroup": "Power",
307 "MetricName": "C2_Pkg_Residency"
310 "BriefDescription": "C6 residency percent per package",
311 "MetricExpr": "(cstate_pkg@c6\\-residency@ / msr@tsc@) * 100",
312 "MetricGroup": "Power",
313 "MetricName": "C6_Pkg_Residency"