Function: hypre_BoomerAMGBuildMultipass.extracted.27 | Module: exec | Source: par_multi_interp.c:1575-1663 [...] | Coverage: 0.21% |
---|
Function: hypre_BoomerAMGBuildMultipass.extracted.27 | Module: exec | Source: par_multi_interp.c:1575-1663 [...] | Coverage: 0.21% |
---|
/scratch_na/users/xoserete/qaas_runs/171-172-8218/intel/AMG/build/AMG/AMG/parcsr_ls/par_multi_interp.c: 1575 - 1663 |
-------------------------------------------------------------------------------- |
1575: #pragma omp parallel private(thread_start,thread_stop,my_thread_num,num_threads,k,k1,i,i1,j,j1,sum_C,sum_N,j_start,j_end,cnt,tmp_marker,tmp_marker_offd,cnt_offd,diagonal,alfa) |
[...] |
1585: if (n_fine) |
1586: { tmp_marker = hypre_CTAlloc(HYPRE_Int,n_fine); } |
1587: tmp_marker_offd = NULL; |
1588: if (num_cols_offd) |
1589: { tmp_marker_offd = hypre_CTAlloc(HYPRE_Int,num_cols_offd); } |
1590: for (i=0; i < n_fine; i++) |
1591: { tmp_marker[i] = -1; } |
1592: for (i=0; i < num_cols_offd; i++) |
1593: { tmp_marker_offd[i] = -1; } |
1594: |
1595: /* Compute this thread's range of pass_length */ |
1596: my_thread_num = hypre_GetThreadNum(); |
1597: num_threads = hypre_NumActiveThreads(); |
1598: thread_start = pass_pointer[1] + (pass_length/num_threads)*my_thread_num; |
1599: if (my_thread_num == num_threads-1) |
[...] |
1605: for (i=thread_start; i < thread_stop; i++) |
1606: { |
1607: i1 = pass_array[i]; |
1608: sum_C = 0; |
1609: sum_N = 0; |
1610: j_start = P_diag_start[i1]; |
1611: j_end = j_start+P_diag_i[i1+1]-P_diag_i[i1]; |
1612: for (j=j_start; j < j_end; j++) |
1613: { |
1614: k1 = P_diag_pass[1][j]; |
1615: tmp_marker[C_array[k1]] = i1; |
1616: } |
1617: cnt = P_diag_i[i1]; |
1618: for (j=A_diag_i[i1]+1; j < A_diag_i[i1+1]; j++) |
1619: { |
1620: j1 = A_diag_j[j]; |
1621: if (CF_marker[j1] != -3 && |
1622: (num_functions == 1 || dof_func[i1] == dof_func[j1])) |
1623: sum_N += A_diag_data[j]; |
1624: if (j1 != -1 && tmp_marker[j1] == i1) |
1625: { |
1626: P_diag_data[cnt] = A_diag_data[j]; |
1627: P_diag_j[cnt++] = fine_to_coarse[j1]; |
1628: sum_C += A_diag_data[j]; |
1629: } |
1630: } |
1631: j_start = P_offd_start[i1]; |
1632: j_end = j_start+P_offd_i[i1+1]-P_offd_i[i1]; |
1633: for (j=j_start; j < j_end; j++) |
1634: { |
1635: k1 = P_offd_pass[1][j]; |
1636: tmp_marker_offd[C_array_offd[k1]] = i1; |
1637: } |
1638: cnt_offd = P_offd_i[i1]; |
1639: for (j=A_offd_i[i1]; j < A_offd_i[i1+1]; j++) |
1640: { |
1641: if (col_offd_S_to_A) |
1642: j1 = map_A_to_S[A_offd_j[j]]; |
1643: else |
1644: j1 = A_offd_j[j]; |
1645: if (CF_marker_offd[j1] != -3 && |
1646: (num_functions == 1 || dof_func[i1] == dof_func_offd[j1])) |
1647: sum_N += A_offd_data[j]; |
1648: if (j1 != -1 && tmp_marker_offd[j1] == i1) |
1649: { |
1650: P_offd_data[cnt_offd] = A_offd_data[j]; |
1651: P_offd_j[cnt_offd++] = map_S_to_new[j1]; |
1652: sum_C += A_offd_data[j]; |
1653: } |
1654: } |
1655: diagonal = A_diag_data[A_diag_i[i1]]; |
1656: if (sum_C*diagonal) alfa = -sum_N/(sum_C*diagonal); |
1657: for (j=P_diag_i[i1]; j < cnt; j++) |
1658: P_diag_data[j] *= alfa; |
1659: for (j=P_offd_i[i1]; j < cnt_offd; j++) |
1660: P_offd_data[j] *= alfa; |
1661: } |
1662: hypre_TFree(tmp_marker); |
1663: hypre_TFree(tmp_marker_offd); |
0x442c00 PUSH %RBP |
0x442c01 MOV %RSP,%RBP |
0x442c04 PUSH %R15 |
0x442c06 PUSH %R14 |
0x442c08 PUSH %R13 |
0x442c0a PUSH %R12 |
0x442c0c PUSH %RBX |
0x442c0d SUB $0xe8,%RSP |
0x442c14 MOV %R9,-0xd0(%RBP) |
0x442c1b MOV %R8,-0x90(%RBP) |
0x442c22 MOV %RCX,-0xa0(%RBP) |
0x442c29 MOV %RDX,-0x40(%RBP) |
0x442c2d MOV 0xe8(%RBP),%RAX |
0x442c34 MOV %RAX,-0x50(%RBP) |
0x442c38 MOV 0xe0(%RBP),%RAX |
0x442c3f MOV %RAX,-0x108(%RBP) |
0x442c46 MOV 0xd8(%RBP),%RDI |
0x442c4d MOV 0xd0(%RBP),%RAX |
0x442c54 MOV %RAX,-0x110(%RBP) |
0x442c5b MOV 0xc8(%RBP),%RAX |
0x442c62 MOV %RAX,-0x100(%RBP) |
0x442c69 MOV 0xc0(%RBP),%RAX |
0x442c70 MOV %RAX,-0xc8(%RBP) |
0x442c77 MOV 0xb8(%RBP),%RAX |
0x442c7e MOV %RAX,-0xc0(%RBP) |
0x442c85 MOV 0xb0(%RBP),%RAX |
0x442c8c MOV %RAX,-0xe8(%RBP) |
0x442c93 MOV 0xa8(%RBP),%RAX |
0x442c9a MOV %RAX,-0xe0(%RBP) |
0x442ca1 MOV 0xa0(%RBP),%RAX |
0x442ca8 MOV %RAX,-0x38(%RBP) |
0x442cac MOV 0x98(%RBP),%RAX |
0x442cb3 MOV %RAX,-0xd8(%RBP) |
0x442cba MOV 0x90(%RBP),%RBX |
0x442cc1 MOV 0x88(%RBP),%R15 |
0x442cc8 MOV 0x80(%RBP),%RAX |
0x442ccf MOV %RAX,-0xf0(%RBP) |
0x442cd6 MOV 0x78(%RBP),%R12 |
0x442cda MOV 0x70(%RBP),%RAX |
0x442cde MOV %RAX,-0xf8(%RBP) |
0x442ce5 MOV 0x68(%RBP),%RAX |
0x442ce9 MOV %RAX,-0x60(%RBP) |
0x442ced MOV 0x60(%RBP),%RAX |
0x442cf1 MOV %RAX,-0x48(%RBP) |
0x442cf5 MOV 0x58(%RBP),%RAX |
0x442cf9 MOV %RAX,-0xb8(%RBP) |
0x442d00 MOV 0x50(%RBP),%RAX |
0x442d04 MOV %RAX,-0x58(%RBP) |
0x442d08 MOV 0x48(%RBP),%RAX |
0x442d0c MOV %RAX,-0x88(%RBP) |
0x442d13 MOV 0x40(%RBP),%RCX |
0x442d17 MOV 0x38(%RBP),%RAX |
0x442d1b MOV %RAX,-0xb0(%RBP) |
0x442d22 MOV 0x30(%RBP),%RAX |
0x442d26 MOV %RAX,-0x80(%RBP) |
0x442d2a MOV 0x28(%RBP),%RAX |
0x442d2e MOV %RAX,-0x98(%RBP) |
0x442d35 MOV 0x20(%RBP),%RAX |
0x442d39 MOV %RAX,-0x78(%RBP) |
0x442d3d MOV 0x18(%RBP),%RAX |
0x442d41 MOV %RAX,-0x68(%RBP) |
0x442d45 MOV 0x10(%RBP),%RAX |
0x442d49 MOV %RAX,-0x70(%RBP) |
0x442d4d TEST %RDI,%RDI |
0x442d50 MOV %RCX,-0x30(%RBP) |
0x442d54 MOV %RDI,-0xa8(%RBP) |
0x442d5b JE 442df5 |
0x442d61 MOV $0x8,%ESI |
0x442d66 CALL 4e6d80 <hypre_CAlloc> |
0x442d6b MOV -0x30(%RBP),%RCX |
0x442d6f MOV %RAX,%R13 |
0x442d72 TEST %RCX,%RCX |
0x442d75 JE 442e01 |
0x442d7b MOV $0x8,%ESI |
0x442d80 MOV %RCX,%RDI |
0x442d83 CALL 4e6d80 <hypre_CAlloc> |
0x442d88 MOV %RAX,%R14 |
0x442d8b MOV -0xa8(%RBP),%RDX |
0x442d92 TEST %RDX,%RDX |
0x442d95 JLE 442da8 |
0x442d97 SAL $0x3,%RDX |
0x442d9b MOV %R13,%RDI |
0x442d9e MOV $0xff,%ESI |
0x442da3 CALL 4efe80 <_intel_fast_memset> |
0x442da8 MOV -0x30(%RBP),%RDX |
0x442dac TEST %RDX,%RDX |
0x442daf JLE 442dc2 |
0x442db1 SAL $0x3,%RDX |
0x442db5 MOV %R14,%RDI |
0x442db8 MOV $0xff,%ESI |
0x442dbd CALL 4efe80 <_intel_fast_memset> |
0x442dc2 CALL 4e8ab0 <hypre_GetThreadNum> |
0x442dc7 MOV %RAX,-0x30(%RBP) |
0x442dcb CALL 4e8aa0 <hypre_NumActiveThreads> |
0x442dd0 MOV %RAX,%RCX |
0x442dd3 MOV -0x38(%RBP),%RAX |
0x442dd7 MOV 0x8(%RAX),%RDI |
0x442ddb MOV -0x50(%RBP),%R8 |
0x442ddf MOV %R8,%RAX |
0x442de2 OR %RCX,%RAX |
0x442de5 SHR $0x20,%RAX |
0x442de9 JE 442e12 |
0x442deb MOV %R8,%RAX |
0x442dee CQTO |
0x442df0 IDIV %RCX |
0x442df3 JMP 442e19 |
0x442df5 XOR %R13D,%R13D |
0x442df8 TEST %RCX,%RCX |
0x442dfb JNE 442d7b |
0x442e01 XOR %R14D,%R14D |
0x442e04 MOV -0xa8(%RBP),%RDX |
0x442e0b TEST %RDX,%RDX |
0x442e0e JG 442d97 |
0x442e10 JMP 442da8 |
0x442e12 MOV %R8D,%EAX |
0x442e15 XOR %EDX,%EDX |
0x442e17 DIV %ECX |
0x442e19 MOV -0x40(%RBP),%R10 |
0x442e1d MOV %RAX,%RDX |
0x442e20 MOV -0x30(%RBP),%R9 |
0x442e24 IMUL %R9,%RDX |
0x442e28 DEC %RCX |
0x442e2b LEA 0x1(%R9),%RSI |
0x442e2f IMUL %RAX,%RSI |
0x442e33 CMP %RCX,%R9 |
0x442e36 CMOVE %R8,%RSI |
0x442e3a MOV %RSI,-0x38(%RBP) |
0x442e3e CMP %RSI,%RDX |
0x442e41 JGE 443367 |
0x442e47 MOV -0x38(%RBP),%RAX |
0x442e4b ADD %RDI,%RAX |
0x442e4e MOV %RAX,-0x38(%RBP) |
0x442e52 ADD %RDI,%RDX |
0x442e55 VXORPD %XMM0,%XMM0,%XMM0 |
0x442e59 VMOVDDUP 0xbc35f(%RIP),%XMM1 |
0x442e61 JMP 442e81 |
0x442e63 NOPW %CS:(%RAX,%RAX,1) |
(955) 0x442e70 MOV -0x50(%RBP),%RDX |
(955) 0x442e74 INC %RDX |
(955) 0x442e77 CMP -0x38(%RBP),%RDX |
(955) 0x442e7b JGE 443367 |
(955) 0x442e81 MOV -0xd8(%RBP),%RAX |
(955) 0x442e88 MOV %RDX,-0x50(%RBP) |
(955) 0x442e8c MOV (%RAX,%RDX,8),%RAX |
(955) 0x442e90 MOV -0xe0(%RBP),%RCX |
(955) 0x442e97 MOV (%RCX,%RAX,8),%RDI |
(955) 0x442e9b MOV -0x58(%RBP),%RCX |
(955) 0x442e9f MOV (%RCX,%RAX,8),%RSI |
(955) 0x442ea3 MOV 0x8(%RCX,%RAX,8),%RCX |
(955) 0x442ea8 LEA (%RCX,%RDI,1),%RDX |
(955) 0x442eac SUB %RSI,%RDX |
(955) 0x442eaf CMP %RDX,%RDI |
(955) 0x442eb2 JGE 442f85 |
(955) 0x442eb8 MOV -0xc0(%RBP),%RDX |
(955) 0x442ebf MOV 0x8(%RDX),%R8 |
(955) 0x442ec3 SUB %RSI,%RCX |
(955) 0x442ec6 CMP $0x8,%RCX |
(955) 0x442eca JB 442f50 |
(955) 0x442ed0 MOV %RCX,%R9 |
(955) 0x442ed3 SHR $0x3,%R9 |
(955) 0x442ed7 LEA 0x38(%R8,%RDI,8),%R10 |
(955) 0x442edc NOPL (%RAX) |
(965) 0x442ee0 MOV -0x38(%R10),%RDX |
(965) 0x442ee4 MOV (%R15,%RDX,8),%RDX |
(965) 0x442ee8 MOV %RAX,(%R13,%RDX,8) |
(965) 0x442eed MOV -0x30(%R10),%RDX |
(965) 0x442ef1 MOV (%R15,%RDX,8),%RDX |
(965) 0x442ef5 MOV %RAX,(%R13,%RDX,8) |
(965) 0x442efa MOV -0x28(%R10),%RDX |
(965) 0x442efe MOV (%R15,%RDX,8),%RDX |
(965) 0x442f02 MOV %RAX,(%R13,%RDX,8) |
(965) 0x442f07 MOV -0x20(%R10),%RDX |
(965) 0x442f0b MOV (%R15,%RDX,8),%RDX |
(965) 0x442f0f MOV %RAX,(%R13,%RDX,8) |
(965) 0x442f14 MOV -0x18(%R10),%RDX |
(965) 0x442f18 MOV (%R15,%RDX,8),%RDX |
(965) 0x442f1c MOV %RAX,(%R13,%RDX,8) |
(965) 0x442f21 MOV -0x10(%R10),%RDX |
(965) 0x442f25 MOV (%R15,%RDX,8),%RDX |
(965) 0x442f29 MOV %RAX,(%R13,%RDX,8) |
(965) 0x442f2e MOV -0x8(%R10),%RDX |
(965) 0x442f32 MOV (%R15,%RDX,8),%RDX |
(965) 0x442f36 MOV %RAX,(%R13,%RDX,8) |
(965) 0x442f3b MOV (%R10),%RDX |
(965) 0x442f3e MOV (%R15,%RDX,8),%RDX |
(965) 0x442f42 MOV %RAX,(%R13,%RDX,8) |
(965) 0x442f47 ADD $0x40,%R10 |
(965) 0x442f4b DEC %R9 |
(965) 0x442f4e JNE 442ee0 |
(955) 0x442f50 MOV %RCX,%R9 |
(955) 0x442f53 AND $-0x8,%R9 |
(955) 0x442f57 CMP %RCX,%R9 |
(955) 0x442f5a MOV -0x40(%RBP),%R10 |
(955) 0x442f5e JAE 442f85 |
(955) 0x442f60 LEA (%R8,%RDI,8),%RDI |
(955) 0x442f64 NOPW %CS:(%RAX,%RAX,1) |
(964) 0x442f70 MOV (%RDI,%R9,8),%RDX |
(964) 0x442f74 MOV (%R15,%RDX,8),%RDX |
(964) 0x442f78 MOV %RAX,(%R13,%RDX,8) |
(964) 0x442f7d INC %R9 |
(964) 0x442f80 CMP %R9,%RCX |
(964) 0x442f83 JNE 442f70 |
(955) 0x442f85 MOV -0x58(%RBP),%RCX |
(955) 0x442f89 MOV (%RCX,%RAX,8),%RDX |
(955) 0x442f8d MOV -0x68(%RBP),%RCX |
(955) 0x442f91 MOV (%RCX,%RAX,8),%RDI |
(955) 0x442f95 MOV 0x8(%RCX,%RAX,8),%R8 |
(955) 0x442f9a INC %RDI |
(955) 0x442f9d VXORPD %XMM4,%XMM4,%XMM4 |
(955) 0x442fa1 CMP %R8,%RDI |
(955) 0x442fa4 MOV %RDX,-0x30(%RBP) |
(955) 0x442fa8 VXORPD %XMM3,%XMM3,%XMM3 |
(955) 0x442fac JGE 443060 |
(955) 0x442fb2 MOV -0xb8(%RBP),%RCX |
(955) 0x442fb9 MOV -0x78(%RBP),%RSI |
(955) 0x442fbd JMP 442fcc |
0x442fbf NOP |
(963) 0x442fc0 INC %RDI |
(963) 0x442fc3 CMP %R8,%RDI |
(963) 0x442fc6 JGE 443060 |
(963) 0x442fcc MOV (%RSI,%RDI,8),%R9 |
(963) 0x442fd0 CMPQ $-0x3,(%R10,%R9,8) |
(963) 0x442fd5 JE 442fff |
(963) 0x442fd7 CMPQ $0x1,-0xa0(%RBP) |
(963) 0x442fdf JE 442ff6 |
(963) 0x442fe1 MOV -0x90(%RBP),%RSI |
(963) 0x442fe8 MOV (%RSI,%RAX,8),%RDX |
(963) 0x442fec CMP (%RSI,%R9,8),%RDX |
(963) 0x442ff0 MOV -0x78(%RBP),%RSI |
(963) 0x442ff4 JNE 442fff |
(963) 0x442ff6 MOV -0x70(%RBP),%RDX |
(963) 0x442ffa VADDSD (%RDX,%RDI,8),%XMM3,%XMM3 |
(963) 0x442fff CMP $-0x1,%R9 |
(963) 0x443003 JE 442fc0 |
(963) 0x443005 CMP %RAX,(%R13,%R9,8) |
(963) 0x44300a JNE 442fc0 |
(963) 0x44300c MOV -0x70(%RBP),%R8 |
(963) 0x443010 VMOVSD (%R8,%RDI,8),%XMM5 |
(963) 0x443016 MOV -0x88(%RBP),%RDX |
(963) 0x44301d MOV -0x30(%RBP),%R11 |
(963) 0x443021 VMOVSD %XMM5,(%RDX,%R11,8) |
(963) 0x443027 MOV -0x108(%RBP),%RDX |
(963) 0x44302e MOV (%RDX,%R9,8),%RDX |
(963) 0x443032 MOV %RDX,(%RCX,%R11,8) |
(963) 0x443036 INC %R11 |
(963) 0x443039 MOV %R11,-0x30(%RBP) |
(963) 0x44303d VADDSD (%R8,%RDI,8),%XMM4,%XMM4 |
(963) 0x443043 MOV -0x68(%RBP),%RDX |
(963) 0x443047 MOV 0x8(%RDX,%RAX,8),%R8 |
(963) 0x44304c JMP 442fc0 |
0x443051 NOPW %CS:(%RAX,%RAX,1) |
(955) 0x443060 MOV -0xe8(%RBP),%RDX |
(955) 0x443067 MOV (%RDX,%RAX,8),%R8 |
(955) 0x44306b MOV -0x60(%RBP),%RCX |
(955) 0x44306f MOV (%RCX,%RAX,8),%RSI |
(955) 0x443073 MOV 0x8(%RCX,%RAX,8),%RDI |
(955) 0x443078 LEA (%RDI,%R8,1),%RDX |
(955) 0x44307c SUB %RSI,%RDX |
(955) 0x44307f CMP %RDX,%R8 |
(955) 0x443082 JGE 443144 |
(955) 0x443088 MOV -0xc8(%RBP),%RDX |
(955) 0x44308f MOV 0x8(%RDX),%R9 |
(955) 0x443093 SUB %RSI,%RDI |
(955) 0x443096 CMP $0x8,%RDI |
(955) 0x44309a JB 443118 |
(955) 0x4430a0 MOV %RDI,%R10 |
(955) 0x4430a3 SHR $0x3,%R10 |
(955) 0x4430a7 LEA 0x38(%R9,%R8,8),%R11 |
(955) 0x4430ac NOPL (%RAX) |
(962) 0x4430b0 MOV -0x38(%R11),%RDX |
(962) 0x4430b4 MOV (%RBX,%RDX,8),%RDX |
(962) 0x4430b8 MOV %RAX,(%R14,%RDX,8) |
(962) 0x4430bc MOV -0x30(%R11),%RDX |
(962) 0x4430c0 MOV (%RBX,%RDX,8),%RDX |
(962) 0x4430c4 MOV %RAX,(%R14,%RDX,8) |
(962) 0x4430c8 MOV -0x28(%R11),%RDX |
(962) 0x4430cc MOV (%RBX,%RDX,8),%RDX |
(962) 0x4430d0 MOV %RAX,(%R14,%RDX,8) |
(962) 0x4430d4 MOV -0x20(%R11),%RDX |
(962) 0x4430d8 MOV (%RBX,%RDX,8),%RDX |
(962) 0x4430dc MOV %RAX,(%R14,%RDX,8) |
(962) 0x4430e0 MOV -0x18(%R11),%RDX |
(962) 0x4430e4 MOV (%RBX,%RDX,8),%RDX |
(962) 0x4430e8 MOV %RAX,(%R14,%RDX,8) |
(962) 0x4430ec MOV -0x10(%R11),%RDX |
(962) 0x4430f0 MOV (%RBX,%RDX,8),%RDX |
(962) 0x4430f4 MOV %RAX,(%R14,%RDX,8) |
(962) 0x4430f8 MOV -0x8(%R11),%RDX |
(962) 0x4430fc MOV (%RBX,%RDX,8),%RDX |
(962) 0x443100 MOV %RAX,(%R14,%RDX,8) |
(962) 0x443104 MOV (%R11),%RDX |
(962) 0x443107 MOV (%RBX,%RDX,8),%RDX |
(962) 0x44310b MOV %RAX,(%R14,%RDX,8) |
(962) 0x44310f ADD $0x40,%R11 |
(962) 0x443113 DEC %R10 |
(962) 0x443116 JNE 4430b0 |
(955) 0x443118 MOV %RDI,%R10 |
(955) 0x44311b AND $-0x8,%R10 |
(955) 0x44311f CMP %RDI,%R10 |
(955) 0x443122 JAE 443144 |
(955) 0x443124 LEA (%R9,%R8,8),%R8 |
(955) 0x443128 NOPL (%RAX,%RAX,1) |
(961) 0x443130 MOV (%R8,%R10,8),%RDX |
(961) 0x443134 MOV (%RBX,%RDX,8),%RDX |
(961) 0x443138 MOV %RAX,(%R14,%RDX,8) |
(961) 0x44313c INC %R10 |
(961) 0x44313f CMP %R10,%RDI |
(961) 0x443142 JNE 443130 |
(955) 0x443144 MOV -0x60(%RBP),%RCX |
(955) 0x443148 MOV (%RCX,%RAX,8),%RDI |
(955) 0x44314c MOV -0x80(%RBP),%RCX |
(955) 0x443150 MOV (%RCX,%RAX,8),%R8 |
(955) 0x443154 MOV 0x8(%RCX,%RAX,8),%R10 |
(955) 0x443159 CMP %R10,%R8 |
(955) 0x44315c JGE 443240 |
(955) 0x443162 MOV -0xb0(%RBP),%RCX |
(955) 0x443169 LEA (%RCX,%R8,8),%R9 |
(955) 0x44316d MOV -0xd0(%RBP),%RSI |
(955) 0x443174 JMP 443190 |
0x443176 NOPW %CS:(%RAX,%RAX,1) |
(960) 0x443180 INC %R8 |
(960) 0x443183 ADD $0x8,%R9 |
(960) 0x443187 CMP %R10,%R8 |
(960) 0x44318a JGE 443240 |
(960) 0x443190 MOV %R9,%RDX |
(960) 0x443193 TEST %RSI,%RSI |
(960) 0x443196 JE 4431a6 |
(960) 0x443198 MOV (%R9),%RDX |
(960) 0x44319b MOV -0x110(%RBP),%R11 |
(960) 0x4431a2 LEA (%R11,%RDX,8),%RDX |
(960) 0x4431a6 MOV (%RDX),%R11 |
(960) 0x4431a9 CMPQ $-0x3,(%R12,%R11,8) |
(960) 0x4431ae JE 4431e5 |
(960) 0x4431b0 CMPQ $0x1,-0xa0(%RBP) |
(960) 0x4431b8 JE 4431d8 |
(960) 0x4431ba MOV -0x90(%RBP),%RDX |
(960) 0x4431c1 MOV (%RDX,%RAX,8),%RDX |
(960) 0x4431c5 MOV %R12,%RCX |
(960) 0x4431c8 MOV -0xf0(%RBP),%R12 |
(960) 0x4431cf CMP (%R12,%R11,8),%RDX |
(960) 0x4431d3 MOV %RCX,%R12 |
(960) 0x4431d6 JNE 4431e5 |
(960) 0x4431d8 MOV -0x98(%RBP),%RCX |
(960) 0x4431df VADDSD (%RCX,%R8,8),%XMM3,%XMM3 |
(960) 0x4431e5 CMP $-0x1,%R11 |
(960) 0x4431e9 JE 443180 |
(960) 0x4431eb CMP %RAX,(%R14,%R11,8) |
(960) 0x4431ef JNE 443180 |
(960) 0x4431f1 MOV -0x98(%RBP),%R10 |
(960) 0x4431f8 VMOVSD (%R10,%R8,8),%XMM5 |
(960) 0x4431fe MOV -0x48(%RBP),%RCX |
(960) 0x443202 VMOVSD %XMM5,(%RCX,%RDI,8) |
(960) 0x443207 MOV -0x100(%RBP),%RDX |
(960) 0x44320e MOV (%RDX,%R11,8),%RDX |
(960) 0x443212 MOV -0xf8(%RBP),%RCX |
(960) 0x443219 MOV %RDX,(%RCX,%RDI,8) |
(960) 0x44321d INC %RDI |
(960) 0x443220 VADDSD (%R10,%R8,8),%XMM4,%XMM4 |
(960) 0x443226 MOV -0x80(%RBP),%RCX |
(960) 0x44322a MOV 0x8(%RCX,%RAX,8),%R10 |
(960) 0x44322f JMP 443180 |
0x443234 NOPW %CS:(%RAX,%RAX,1) |
(955) 0x443240 MOV -0x68(%RBP),%RCX |
(955) 0x443244 MOV (%RCX,%RAX,8),%RDX |
(955) 0x443248 MOV -0x70(%RBP),%RCX |
(955) 0x44324c VMULSD (%RCX,%RDX,8),%XMM4,%XMM4 |
(955) 0x443251 VUCOMISD %XMM0,%XMM4 |
(955) 0x443255 JE 44325f |
(955) 0x443257 VXORPD %XMM1,%XMM3,%XMM2 |
(955) 0x44325b VDIVSD %XMM4,%XMM2,%XMM2 |
(955) 0x44325f MOV -0x58(%RBP),%RCX |
(955) 0x443263 MOV (%RCX,%RAX,8),%R9 |
(955) 0x443267 MOV -0x30(%RBP),%RSI |
(955) 0x44326b MOV %RSI,%R10 |
(955) 0x44326e SUB %R9,%R10 |
(955) 0x443271 MOV -0x48(%RBP),%RDX |
(955) 0x443275 MOV -0x88(%RBP),%RCX |
(955) 0x44327c JLE 4432e4 |
(955) 0x44327e MOV %R10,%R8 |
(955) 0x443281 AND $-0x4,%R8 |
(955) 0x443285 JE 4432c2 |
(955) 0x443287 LEA -0x1(%R8),%R11 |
(955) 0x44328b VBROADCASTSD %XMM2,%YMM3 |
(955) 0x443290 LEA (%RCX,%R9,8),%RSI |
(955) 0x443294 XOR %EDX,%EDX |
(955) 0x443296 NOPW %CS:(%RAX,%RAX,1) |
(959) 0x4432a0 VMULPD (%RSI,%RDX,8),%YMM3,%YMM4 |
(959) 0x4432a5 VMOVUPD %YMM4,(%RSI,%RDX,8) |
(959) 0x4432aa ADD $0x4,%RDX |
(959) 0x4432ae CMP %R11,%RDX |
(959) 0x4432b1 JBE 4432a0 |
(955) 0x4432b3 CMP %R8,%R10 |
(955) 0x4432b6 MOV -0x48(%RBP),%RDX |
(955) 0x4432ba MOV -0x30(%RBP),%RSI |
(955) 0x4432be JNE 4432c5 |
(955) 0x4432c0 JMP 4432e4 |
(955) 0x4432c2 XOR %R8D,%R8D |
(955) 0x4432c5 ADD %R9,%R8 |
(955) 0x4432c8 NOPL (%RAX,%RAX,1) |
(958) 0x4432d0 VMULSD (%RCX,%R8,8),%XMM2,%XMM3 |
(958) 0x4432d6 VMOVSD %XMM3,(%RCX,%R8,8) |
(958) 0x4432dc INC %R8 |
(958) 0x4432df CMP %R8,%RSI |
(958) 0x4432e2 JNE 4432d0 |
(955) 0x4432e4 MOV -0x60(%RBP),%RCX |
(955) 0x4432e8 MOV (%RCX,%RAX,8),%RCX |
(955) 0x4432ec MOV %RDI,%R8 |
(955) 0x4432ef SUB %RCX,%R8 |
(955) 0x4432f2 MOV -0x40(%RBP),%R10 |
(955) 0x4432f6 JLE 442e70 |
(955) 0x4432fc MOV %R8,%RAX |
(955) 0x4432ff AND $-0x4,%RAX |
(955) 0x443303 JE 443342 |
(955) 0x443305 LEA -0x1(%RAX),%R9 |
(955) 0x443309 VBROADCASTSD %XMM2,%YMM3 |
(955) 0x44330e LEA (%RDX,%RCX,8),%RSI |
(955) 0x443312 XOR %EDX,%EDX |
(955) 0x443314 NOPW %CS:(%RAX,%RAX,1) |
(957) 0x443320 VMULPD (%RSI,%RDX,8),%YMM3,%YMM4 |
(957) 0x443325 VMOVUPD %YMM4,(%RSI,%RDX,8) |
(957) 0x44332a ADD $0x4,%RDX |
(957) 0x44332e CMP %R9,%RDX |
(957) 0x443331 JBE 443320 |
(955) 0x443333 CMP %RAX,%R8 |
(955) 0x443336 MOV -0x48(%RBP),%RDX |
(955) 0x44333a JE 442e70 |
(955) 0x443340 JMP 443344 |
(955) 0x443342 XOR %EAX,%EAX |
(955) 0x443344 ADD %RCX,%RAX |
(955) 0x443347 NOPW (%RAX,%RAX,1) |
(956) 0x443350 VMULSD (%RDX,%RAX,8),%XMM2,%XMM3 |
(956) 0x443355 VMOVSD %XMM3,(%RDX,%RAX,8) |
(956) 0x44335a INC %RAX |
(956) 0x44335d CMP %RAX,%RDI |
(956) 0x443360 JNE 443350 |
(955) 0x443362 JMP 442e70 |
0x443367 MOV %R13,%RDI |
0x44336a VZEROUPPER |
0x44336d CALL 4e6e50 <hypre_Free> |
0x443372 MOV %R14,%RDI |
0x443375 ADD $0xe8,%RSP |
0x44337c POP %RBX |
0x44337d POP %R12 |
0x44337f POP %R13 |
0x443381 POP %R14 |
0x443383 POP %R15 |
0x443385 POP %RBP |
0x443386 JMP 4e6e50 |
0x44338b NOPL (%RAX,%RAX,1) |
Path / |
Source file and lines | par_multi_interp.c:1575-1663 |
Module | exec |
nb instructions | 154 |
nb uops | 169 |
loop length | 703 |
used x86 registers | 15 |
used mmx registers | 0 |
used xmm registers | 2 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 57 |
micro-operation queue | 28.17 cycles |
front end | 28.17 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 7.50 | 8.00 | 15.00 | 15.00 | 22.50 | 7.60 | 7.50 | 22.50 | 22.50 | 22.50 | 7.40 | 15.00 |
cycles | 7.50 | 11.40 | 15.00 | 15.00 | 22.50 | 7.60 | 7.50 | 22.50 | 22.50 | 22.50 | 7.40 | 15.00 |
Cycles executing div or sqrt instructions | 16.00 |
FE+BE cycles | 26.56-26.61 |
Stall cycles | 0.00 |
Front-end | 28.17 |
Dispatch | 22.50 |
DIV/SQRT | 16.00 |
Overall L1 | 28.17 |
all | 1% |
load | 0% |
store | 0% |
mul | 0% |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 3% |
all | 50% |
load | 0% |
store | NA (no store vectorizable/vectorized instructions) |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 100% |
all | 2% |
load | 0% |
store | 0% |
mul | 0% |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 0% |
other | 6% |
all | 12% |
load | 12% |
store | 12% |
mul | 12% |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 11% |
all | 18% |
load | 12% |
store | NA (no store vectorizable/vectorized instructions) |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 25% |
all | 12% |
load | 12% |
store | 12% |
mul | 12% |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 9% |
other | 11% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
SUB $0xe8,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R9,-0xd0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R8,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0xa0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDX,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xe8(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x50(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xe0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x108(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xd8(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xd0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x110(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xc8(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x100(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xc0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xc8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xb8(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xb0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xe8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xa8(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xe0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xa0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x98(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xd8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x90(%RBP),%RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x88(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x80(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xf0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x78(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x70(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xf8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x68(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x60(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x58(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xb8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x50(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xb0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x28(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x18(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x10(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
TEST %RDI,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
MOV %RCX,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDI,-0xa8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JE 442df5 <hypre_BoomerAMGBuildMultipass.extracted.27+0x1f5> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV $0x8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 4e6d80 <hypre_CAlloc> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JE 442e01 <hypre_BoomerAMGBuildMultipass.extracted.27+0x201> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV $0x8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RCX,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 4e6d80 <hypre_CAlloc> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0xa8(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %RDX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 442da8 <hypre_BoomerAMGBuildMultipass.extracted.27+0x1a8> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
SAL $0x3,%RDX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %R13,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV $0xff,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 4efe80 <_intel_fast_memset> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %RDX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 442dc2 <hypre_BoomerAMGBuildMultipass.extracted.27+0x1c2> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
SAL $0x3,%RDX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %R14,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV $0xff,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 4efe80 <_intel_fast_memset> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
CALL 4e8ab0 <hypre_GetThreadNum> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CALL 4e8aa0 <hypre_NumActiveThreads> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x8(%RAX),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R8,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
OR %RCX,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
SHR $0x20,%RAX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
JE 442e12 <hypre_BoomerAMGBuildMultipass.extracted.27+0x212> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R8,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CQTO | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
IDIV %RCX | 5 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 10 |
JMP 442e19 <hypre_BoomerAMGBuildMultipass.extracted.27+0x219> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
XOR %R13D,%R13D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JNE 442d7b <hypre_BoomerAMGBuildMultipass.extracted.27+0x17b> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
XOR %R14D,%R14D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0xa8(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %RDX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JG 442d97 <hypre_BoomerAMGBuildMultipass.extracted.27+0x197> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 442da8 <hypre_BoomerAMGBuildMultipass.extracted.27+0x1a8> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
MOV %R8D,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
DIV %ECX | 4 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 6 |
MOV -0x40(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x30(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %R9,%RDX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
LEA 0x1(%R9),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
IMUL %RAX,%RSI | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
CMP %RCX,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMOVE %R8,%RSI | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %RSI,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CMP %RSI,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 443367 <hypre_BoomerAMGBuildMultipass.extracted.27+0x767> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RDI,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RAX,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
ADD %RDI,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VXORPD %XMM0,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVDDUP 0xbc35f(%RIP),%XMM1 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 442e81 <hypre_BoomerAMGBuildMultipass.extracted.27+0x281> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R13,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 4e6e50 <hypre_Free> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %R14,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
ADD $0xe8,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
JMP 4e6e50 <hypre_Free> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
Source file and lines | par_multi_interp.c:1575-1663 |
Module | exec |
nb instructions | 154 |
nb uops | 169 |
loop length | 703 |
used x86 registers | 15 |
used mmx registers | 0 |
used xmm registers | 2 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 57 |
micro-operation queue | 28.17 cycles |
front end | 28.17 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 7.50 | 8.00 | 15.00 | 15.00 | 22.50 | 7.60 | 7.50 | 22.50 | 22.50 | 22.50 | 7.40 | 15.00 |
cycles | 7.50 | 11.40 | 15.00 | 15.00 | 22.50 | 7.60 | 7.50 | 22.50 | 22.50 | 22.50 | 7.40 | 15.00 |
Cycles executing div or sqrt instructions | 16.00 |
FE+BE cycles | 26.56-26.61 |
Stall cycles | 0.00 |
Front-end | 28.17 |
Dispatch | 22.50 |
DIV/SQRT | 16.00 |
Overall L1 | 28.17 |
all | 1% |
load | 0% |
store | 0% |
mul | 0% |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 3% |
all | 50% |
load | 0% |
store | NA (no store vectorizable/vectorized instructions) |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 100% |
all | 2% |
load | 0% |
store | 0% |
mul | 0% |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 0% |
other | 6% |
all | 12% |
load | 12% |
store | 12% |
mul | 12% |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 11% |
all | 18% |
load | 12% |
store | NA (no store vectorizable/vectorized instructions) |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 25% |
all | 12% |
load | 12% |
store | 12% |
mul | 12% |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 9% |
other | 11% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
SUB $0xe8,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R9,-0xd0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R8,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0xa0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDX,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xe8(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x50(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xe0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x108(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xd8(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xd0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x110(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xc8(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x100(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xc0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xc8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xb8(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xb0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xe8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xa8(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xe0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xa0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x98(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xd8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x90(%RBP),%RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x88(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x80(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xf0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x78(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x70(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xf8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x68(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x60(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x58(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xb8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x50(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0xb0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x28(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x18(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x10(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
TEST %RDI,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
MOV %RCX,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDI,-0xa8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JE 442df5 <hypre_BoomerAMGBuildMultipass.extracted.27+0x1f5> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV $0x8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 4e6d80 <hypre_CAlloc> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JE 442e01 <hypre_BoomerAMGBuildMultipass.extracted.27+0x201> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV $0x8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RCX,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 4e6d80 <hypre_CAlloc> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0xa8(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %RDX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 442da8 <hypre_BoomerAMGBuildMultipass.extracted.27+0x1a8> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
SAL $0x3,%RDX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %R13,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV $0xff,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 4efe80 <_intel_fast_memset> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %RDX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 442dc2 <hypre_BoomerAMGBuildMultipass.extracted.27+0x1c2> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
SAL $0x3,%RDX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %R14,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV $0xff,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 4efe80 <_intel_fast_memset> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
CALL 4e8ab0 <hypre_GetThreadNum> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CALL 4e8aa0 <hypre_NumActiveThreads> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x8(%RAX),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R8,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
OR %RCX,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
SHR $0x20,%RAX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
JE 442e12 <hypre_BoomerAMGBuildMultipass.extracted.27+0x212> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R8,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CQTO | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
IDIV %RCX | 5 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 10 |
JMP 442e19 <hypre_BoomerAMGBuildMultipass.extracted.27+0x219> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
XOR %R13D,%R13D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JNE 442d7b <hypre_BoomerAMGBuildMultipass.extracted.27+0x17b> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
XOR %R14D,%R14D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0xa8(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %RDX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JG 442d97 <hypre_BoomerAMGBuildMultipass.extracted.27+0x197> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 442da8 <hypre_BoomerAMGBuildMultipass.extracted.27+0x1a8> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
MOV %R8D,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
DIV %ECX | 4 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 6 |
MOV -0x40(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x30(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %R9,%RDX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
LEA 0x1(%R9),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
IMUL %RAX,%RSI | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
CMP %RCX,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMOVE %R8,%RSI | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %RSI,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CMP %RSI,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 443367 <hypre_BoomerAMGBuildMultipass.extracted.27+0x767> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RDI,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RAX,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
ADD %RDI,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VXORPD %XMM0,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVDDUP 0xbc35f(%RIP),%XMM1 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 442e81 <hypre_BoomerAMGBuildMultipass.extracted.27+0x281> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R13,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 4e6e50 <hypre_Free> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %R14,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
ADD $0xe8,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
JMP 4e6e50 <hypre_Free> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
Name | Coverage (%) | Time (s) |
---|---|---|
▼hypre_BoomerAMGBuildMultipass.extracted.27– | 0.21 | 0.05 |
▼Loop 955 - par_multi_interp.c:1585-1660 - exec– | 0.03 | 0.01 |
○Loop 963 - par_multi_interp.c:1618-1628 - exec | 0.17 | 0.03 |
○Loop 964 - par_multi_interp.c:1612-1615 - exec | 0.01 | 0 |
○Loop 957 - par_multi_interp.c:1659-1660 - exec | 0 | 0 |
○Loop 958 - par_multi_interp.c:1657-1658 - exec | 0 | 0 |
○Loop 960 - par_multi_interp.c:1622-1652 - exec | 0 | 0 |
○Loop 962 - par_multi_interp.c:1633-1636 - exec | 0 | 0 |
○Loop 961 - par_multi_interp.c:1633-1636 - exec | 0 | 0 |
○Loop 965 - par_multi_interp.c:1612-1615 - exec | 0 | 0 |
○Loop 959 - par_multi_interp.c:1657-1658 - exec | 0 | 0 |
○Loop 956 - par_multi_interp.c:1659-1660 - exec | 0 | 0 |