Loop Id: 200 | Module: libIJ_mv.so | Source: IJMatrix_parcsr.c:3262-3484 [...] | Coverage: 0.07% |
---|
Loop Id: 200 | Module: libIJ_mv.so | Source: IJMatrix_parcsr.c:3262-3484 [...] | Coverage: 0.07% |
---|
0x12980 MOV -0x68(%RBP),%RDI |
0x12984 MOV -0x50(%RBP),%RSI |
0x12988 MOV -0x70(%RBP),%RDX |
0x1298c MOV -0x40(%RBP),%RCX |
0x12990 MOV -0x30(%RBP),%R11 |
0x12994 MOV 0x20(%RBP),%RAX |
0x12998 INC %RDX |
0x1299b CMP %RCX,%RDX |
0x1299e JGE 138c0 |
0x129a4 MOV (%RDI,%RDX,8),%R10 |
0x129a8 MOV (%RSI,%RDX,8),%R14 |
0x129ac MOV %R10,-0x60(%RBP) |
0x129b0 SUB (%RAX),%R10 |
0x129b3 JL 12a80 |
0x129b9 MOV -0x60(%RBP),%R8 |
0x129bd CMP 0x8(%RAX),%R8 |
0x129c1 JGE 12a80 |
0x129c7 CMPQ $0,0x58(%RBP) |
0x129cc MOV %R10,-0x48(%RBP) |
0x129d0 JE 12d40 |
0x129d6 MOV 0x38(%RBP),%RAX |
0x129da MOV (%RAX,%R10,8),%R13 |
0x129de MOV 0x40(%RBP),%RAX |
0x129e2 MOV (%RAX,%R10,8),%RAX |
0x129e6 MOV %RAX,-0x98(%RBP) |
0x129ed MOV 0x50(%RBP),%RAX |
0x129f1 MOV (%RAX,%R10,8),%RCX |
0x129f5 MOV 0x48(%RBP),%RAX |
0x129f9 MOV (%RAX,%R10,8),%R12 |
0x129fd MOV %RCX,-0x80(%RBP) |
0x12a01 MOV %RCX,%RAX |
0x12a04 SUB %R12,%RAX |
0x12a07 MOV %R14,%RDI |
0x12a0a SUB %RAX,%RDI |
0x12a0d MOV %RDX,-0x70(%RBP) |
0x12a11 MOV %R11,-0x30(%RBP) |
0x12a15 JLE 13240 |
0x12a1b MOV $0x8,%ESI |
0x12a20 MOV %RDI,-0x60(%RBP) |
0x12a24 VZEROUPPER |
0x12a27 CALL 3360 <hypre_CAlloc@plt> |
0x12a2c MOV %RAX,-0x88(%RBP) |
0x12a33 MOV $0x8,%ESI |
0x12a38 MOV -0x60(%RBP),%RDI |
0x12a3c CALL 3360 <hypre_CAlloc@plt> |
0x12a41 MOV -0x30(%RBP),%R11 |
0x12a45 MOV %RAX,-0x78(%RBP) |
0x12a49 TEST %R14,%R14 |
0x12a4c JG 13254 |
0x12a52 MOV 0x48(%RBP),%RAX |
0x12a56 MOV -0x48(%RBP),%RCX |
0x12a5a MOV %R12,(%RAX,%RCX,8) |
0x12a5e JMP 13684 |
0x12a80 ADD %R14,%R11 |
0x12a83 CMPB $0,-0x31(%RBP) |
0x12a87 JNE 12998 |
0x12a8d TEST %R14,%R14 |
0x12a90 JLE 12998 |
0x12a96 MOV %R11,-0x30(%RBP) |
0x12a9a MOV %RDX,-0x70(%RBP) |
0x12a9e MOV 0x98(%RBP),%RAX |
0x12aa5 DEC %RAX |
0x12aa8 SHR $0x1,%RAX |
0x12aab MOV %RAX,-0x90(%RBP) |
0x12ab2 DEC %R14 |
0x12ab5 XOR %EDX,%EDX |
0x12ab7 XOR %ECX,%ECX |
0x12ab9 JMP 12ad1 |
(201) 0x12ac0 CMP -0x90(%RBP),%RCX |
(201) 0x12ac7 LEA 0x1(%RCX),%RCX |
(201) 0x12acb JE 12980 |
(201) 0x12ad1 MOV %RDX,%RAX |
(201) 0x12ad4 MOV %RCX,%RDI |
(201) 0x12ad7 SAL $0x4,%RDI |
(201) 0x12adb MOV 0xa0(%RBP),%R8 |
(201) 0x12ae2 MOV 0x8(%R8,%RDI,1),%RSI |
(201) 0x12ae7 ADD %RSI,%RDX |
(201) 0x12aea MOV -0x60(%RBP),%R9 |
(201) 0x12aee CMP %R9,(%R8,%RDI,1) |
(201) 0x12af2 JNE 12ac0 |
(201) 0x12af4 TEST %RSI,%RSI |
(201) 0x12af7 JLE 12ac0 |
(201) 0x12af9 MOV 0xa8(%RBP),%RDI |
(201) 0x12b00 LEA -0x8(%RDI,%RDX,8),%R8 |
(201) 0x12b05 LEA (%RDI,%RAX,8),%RDI |
(201) 0x12b09 CMP %R15,%R8 |
(201) 0x12b0c JB 12bc0 |
(201) 0x12b12 CMP %RDI,%R15 |
(201) 0x12b15 JB 12bc0 |
(201) 0x12b1b XOR %EAX,%EAX |
(201) 0x12b1d JMP 12b4d |
(205) 0x12b40 CMP %R14,%RAX |
(205) 0x12b43 LEA 0x1(%RAX),%RAX |
(205) 0x12b47 JE 12ac0 |
(205) 0x12b4d MOV (%RBX,%RAX,8),%R8 |
(205) 0x12b51 XOR %R9D,%R9D |
(205) 0x12b54 JMP 12b88 |
(206) 0x12b80 INC %R9 |
(206) 0x12b83 CMP %R9,%RSI |
(206) 0x12b86 JE 12b40 |
(206) 0x12b88 CMP %R8,(%RDI,%R9,8) |
(206) 0x12b8c JNE 12b80 |
(206) 0x12b8e MOVQ $-0x1,(%RDI,%R9,8) |
(206) 0x12b96 INCQ (%R15) |
(206) 0x12b99 JMP 12b80 |
(201) 0x12bc0 MOV -0xa8(%RBP),%R8 |
(201) 0x12bc7 LEA (%R8,%RAX,8),%RAX |
(201) 0x12bcb MOV %RSI,%R12 |
(201) 0x12bce SHR $0x2,%R12 |
(201) 0x12bd2 MOV %RSI,%R10 |
(201) 0x12bd5 AND $-0x4,%R10 |
(201) 0x12bd9 XOR %R11D,%R11D |
(201) 0x12bdc JMP 12c0d |
(202) 0x12c00 CMP %R14,%R11 |
(202) 0x12c03 LEA 0x1(%R11),%R11 |
(202) 0x12c07 JE 12ac0 |
(202) 0x12c0d MOV (%RBX,%R11,8),%R13 |
(202) 0x12c11 CMP $0x4,%RSI |
(202) 0x12c15 JAE 12c80 |
(202) 0x12c17 CMP %RSI,%R10 |
(202) 0x12c1a JAE 12c00 |
(202) 0x12c1c MOV %R10,%R8 |
(202) 0x12c1f JMP 12c48 |
(203) 0x12c40 INC %R8 |
(203) 0x12c43 CMP %R8,%RSI |
(203) 0x12c46 JE 12c00 |
(203) 0x12c48 CMP %R13,(%RDI,%R8,8) |
(203) 0x12c4c JNE 12c40 |
(203) 0x12c4e MOVQ $-0x1,(%RDI,%R8,8) |
(203) 0x12c56 INCQ (%R15) |
(203) 0x12c59 JMP 12c40 |
(202) 0x12c80 MOV %R12,%R9 |
(202) 0x12c83 MOV %RAX,%R8 |
(202) 0x12c86 JMP 12ccd |
(204) 0x12cc0 ADD $0x20,%R8 |
(204) 0x12cc4 DEC %R9 |
(204) 0x12cc7 JE 12c17 |
(204) 0x12ccd CMP %R13,-0x18(%R8) |
(204) 0x12cd1 JNE 12d00 |
(204) 0x12cd3 MOVQ $-0x1,-0x18(%R8) |
(204) 0x12cdb INCQ (%R15) |
(204) 0x12cde CMP %R13,-0x10(%R8) |
(204) 0x12ce2 JE 12d06 |
(204) 0x12ce4 CMP %R13,-0x8(%R8) |
(204) 0x12ce8 JNE 12d17 |
(204) 0x12cea MOVQ $-0x1,-0x8(%R8) |
(204) 0x12cf2 INCQ (%R15) |
(204) 0x12cf5 CMP %R13,(%R8) |
(204) 0x12cf8 JNE 12cc0 |
(204) 0x12cfa JMP 12d1c |
(204) 0x12d00 CMP %R13,-0x10(%R8) |
(204) 0x12d04 JNE 12ce4 |
(204) 0x12d06 MOVQ $-0x1,-0x10(%R8) |
(204) 0x12d0e INCQ (%R15) |
(204) 0x12d11 CMP %R13,-0x8(%R8) |
(204) 0x12d15 JE 12cea |
(204) 0x12d17 CMP %R13,(%R8) |
(204) 0x12d1a JNE 12cc0 |
(204) 0x12d1c MOVQ $-0x1,(%R8) |
(204) 0x12d23 INCQ (%R15) |
(204) 0x12d26 JMP 12cc0 |
0x12d40 MOV 0x18(%RBP),%RCX |
0x12d44 MOV 0x38(%RCX),%RAX |
0x12d48 MOV 0x40(%RCX),%RCX |
0x12d4c MOV (%RCX,%R10,8),%R8 |
0x12d50 MOV (%RAX,%R10,8),%RCX |
0x12d54 TEST %R14,%R14 |
0x12d57 JLE 13600 |
0x12d5d MOV %RDX,-0x70(%RBP) |
0x12d61 MOV 0x60(%RBP),%RAX |
0x12d65 MOV 0x8(%RAX,%R10,8),%RAX |
0x12d6a MOV %RAX,-0x98(%RBP) |
0x12d71 MOV 0x78(%RBP),%RAX |
0x12d75 MOV 0x8(%RAX,%R10,8),%RAX |
0x12d7a MOV %RAX,-0x88(%RBP) |
0x12d81 XOR %EDI,%EDI |
0x12d83 MOV %R8,-0x58(%RBP) |
0x12d87 MOV %RCX,-0x90(%RBP) |
0x12d8e MOV %RCX,-0x80(%RBP) |
0x12d92 LEA (%R11,%R14,1),%RAX |
0x12d96 MOV %RAX,-0xa0(%RBP) |
0x12d9d MOV 0x88(%RBP),%RSI |
0x12da4 MOV %R8,%RCX |
0x12da7 JMP 12dd6 |
(207) 0x12dc0 MOV -0x48(%RBP),%R10 |
(207) 0x12dc4 MOV %RCX,%R8 |
(207) 0x12dc7 INC %R11 |
(207) 0x12dca INC %RDI |
(207) 0x12dcd CMP %R14,%RDI |
(207) 0x12dd0 JE 13640 |
(207) 0x12dd6 MOV %R11,-0x30(%RBP) |
(207) 0x12dda MOV (%RBX,%R11,8),%R9 |
(207) 0x12dde CMP 0x28(%RBP),%R9 |
(207) 0x12de2 JL 12f40 |
(207) 0x12de8 CMP 0x30(%RBP),%R9 |
(207) 0x12dec JG 12f40 |
(207) 0x12df2 MOV 0x60(%RBP),%RAX |
(207) 0x12df6 MOV (%RAX,%R10,8),%R10 |
(207) 0x12dfa MOV -0x90(%RBP),%R11 |
(207) 0x12e01 SUB %R10,%R11 |
(207) 0x12e04 JLE 12f00 |
(207) 0x12e0a MOV 0x68(%RBP),%RAX |
(207) 0x12e0e LEA (%RAX,%R10,8),%EAX |
(207) 0x12e12 AND $0x7f,%EAX |
(207) 0x12e15 MOV $0x80,%EDX |
(207) 0x12e1a SUB %EAX,%EDX |
(207) 0x12e1c SHR $0x3,%EDX |
(207) 0x12e1f CMP %RDX,%R11 |
(207) 0x12e22 MOV %RDX,%R8 |
(207) 0x12e25 CMOVB %R11,%R8 |
(207) 0x12e29 TEST %R8,%R8 |
(207) 0x12e2c JE 12e56 |
(207) 0x12e2e MOV %R10,%R13 |
(207) 0x12e31 MOV %R8,%RAX |
(207) 0x12e34 NOPW %CS:(%RAX,%RAX,1) |
(213) 0x12e40 MOV 0x68(%RBP),%R12 |
(213) 0x12e44 CMP %R9,(%R12,%R13,8) |
(213) 0x12e48 JE 130c0 |
(213) 0x12e4e INC %R13 |
(213) 0x12e51 DEC %RAX |
(213) 0x12e54 JNE 12e40 |
(207) 0x12e56 CMP %RDX,%R11 |
(207) 0x12e59 JBE 12f00 |
(207) 0x12e5f SUB %R8,%R11 |
(207) 0x12e62 MOV %R11,%R12 |
(207) 0x12e65 AND $-0x10,%R12 |
(207) 0x12e69 JE 12eca |
(207) 0x12e6b LEA -0x1(%R12),%RSI |
(207) 0x12e70 LEA (%R10,%R8,1),%R13 |
(207) 0x12e74 MOV 0x68(%RBP),%RAX |
(207) 0x12e78 LEA (%RAX,%R13,8),%RAX |
(207) 0x12e7c VPBROADCASTQ %R9,%YMM0 |
(207) 0x12e82 XOR %EDX,%EDX |
(207) 0x12e84 NOPW %CS:(%RAX,%RAX,1) |
(212) 0x12e90 VPCMPEQQ 0x20(%RAX,%RDX,8),%YMM0,%K0 |
(212) 0x12e98 VPCMPEQQ (%RAX,%RDX,8),%YMM0,%K1 |
(212) 0x12e9f VPCMPEQQ 0x60(%RAX,%RDX,8),%YMM0,%K2 |
(212) 0x12ea7 VPCMPEQQ 0x40(%RAX,%RDX,8),%YMM0,%K3 |
(212) 0x12eaf KORB %K2,%K3,%K4 |
(212) 0x12eb3 KORB %K0,%K1,%K5 |
(212) 0x12eb7 KORTESTB %K4,%K5 |
(212) 0x12ebb JNE 13140 |
(212) 0x12ec1 ADD $0x10,%RDX |
(212) 0x12ec5 CMP %RSI,%RDX |
(212) 0x12ec8 JBE 12e90 |
(207) 0x12eca CMP %R11,%R12 |
(207) 0x12ecd MOV 0x88(%RBP),%RSI |
(207) 0x12ed4 JAE 12f00 |
(207) 0x12ed6 ADD %R8,%R10 |
(207) 0x12ed9 ADD %R12,%R10 |
(207) 0x12edc MOV %R10,%R13 |
(207) 0x12edf NOP |
(211) 0x12ee0 MOV 0x68(%RBP),%RAX |
(211) 0x12ee4 CMP %R9,(%RAX,%R13,8) |
(211) 0x12ee8 JE 130c0 |
(211) 0x12eee INC %R13 |
(211) 0x12ef1 CMP %R13,-0x90(%RBP) |
(211) 0x12ef8 JNE 12ee0 |
(207) 0x12efa NOPW (%RAX,%RAX,1) |
(207) 0x12f00 MOV -0x80(%RBP),%RDX |
(207) 0x12f04 CMP -0x98(%RBP),%RDX |
(207) 0x12f0b JGE 13700 |
(207) 0x12f11 MOV 0x68(%RBP),%RAX |
(207) 0x12f15 MOV %R9,(%RAX,%RDX,8) |
(207) 0x12f19 MOV -0x30(%RBP),%R11 |
(207) 0x12f1d MOV 0x10(%RBP),%RAX |
(207) 0x12f21 VMOVQ (%RAX,%R11,8),%XMM0 |
(207) 0x12f27 MOV 0x70(%RBP),%RAX |
(207) 0x12f2b VMOVQ %XMM0,(%RAX,%RDX,8) |
(207) 0x12f30 INC %RDX |
(207) 0x12f33 MOV %RDX,-0x80(%RBP) |
(207) 0x12f37 JMP 12dc0 |
(207) 0x12f40 MOV 0x78(%RBP),%RAX |
(207) 0x12f44 MOV (%RAX,%R10,8),%R10 |
(207) 0x12f48 MOV %R8,%R11 |
(207) 0x12f4b SUB %R10,%R11 |
(207) 0x12f4e JLE 13050 |
(207) 0x12f54 MOV 0x80(%RBP),%RAX |
(207) 0x12f5b LEA (%RAX,%R10,8),%EAX |
(207) 0x12f5f AND $0x7f,%EAX |
(207) 0x12f62 MOV $0x80,%EDX |
(207) 0x12f67 SUB %EAX,%EDX |
(207) 0x12f69 SHR $0x3,%EDX |
(207) 0x12f6c CMP %RDX,%R11 |
(207) 0x12f6f MOV %RDX,%R8 |
(207) 0x12f72 CMOVB %R11,%R8 |
(207) 0x12f76 TEST %R8,%R8 |
(207) 0x12f79 JE 12fa9 |
(207) 0x12f7b MOV %R10,%R13 |
(207) 0x12f7e MOV %R8,%RAX |
(207) 0x12f81 NOPW %CS:(%RAX,%RAX,1) |
(210) 0x12f90 MOV 0x80(%RBP),%R12 |
(210) 0x12f97 CMP %R9,(%R12,%R13,8) |
(210) 0x12f9b JE 13100 |
(210) 0x12fa1 INC %R13 |
(210) 0x12fa4 DEC %RAX |
(210) 0x12fa7 JNE 12f90 |
(207) 0x12fa9 CMP %RDX,%R11 |
(207) 0x12fac JBE 13050 |
(207) 0x12fb2 SUB %R8,%R11 |
(207) 0x12fb5 MOV %R11,%R12 |
(207) 0x12fb8 AND $-0x10,%R12 |
(207) 0x12fbc JE 1301a |
(207) 0x12fbe LEA -0x1(%R12),%RSI |
(207) 0x12fc3 LEA (%R10,%R8,1),%R13 |
(207) 0x12fc7 MOV 0x80(%RBP),%RAX |
(207) 0x12fce LEA (%RAX,%R13,8),%RAX |
(207) 0x12fd2 VPBROADCASTQ %R9,%YMM0 |
(207) 0x12fd8 XOR %EDX,%EDX |
(207) 0x12fda NOPW (%RAX,%RAX,1) |
(209) 0x12fe0 VPCMPEQQ 0x20(%RAX,%RDX,8),%YMM0,%K0 |
(209) 0x12fe8 VPCMPEQQ (%RAX,%RDX,8),%YMM0,%K1 |
(209) 0x12fef VPCMPEQQ 0x60(%RAX,%RDX,8),%YMM0,%K2 |
(209) 0x12ff7 VPCMPEQQ 0x40(%RAX,%RDX,8),%YMM0,%K3 |
(209) 0x12fff KORB %K2,%K3,%K4 |
(209) 0x13003 KORB %K0,%K1,%K5 |
(209) 0x13007 KORTESTB %K4,%K5 |
(209) 0x1300b JNE 131c0 |
(209) 0x13011 ADD $0x10,%RDX |
(209) 0x13015 CMP %RSI,%RDX |
(209) 0x13018 JBE 12fe0 |
(207) 0x1301a CMP %R11,%R12 |
(207) 0x1301d MOV 0x88(%RBP),%RSI |
(207) 0x13024 JAE 13050 |
(207) 0x13026 ADD %R8,%R10 |
(207) 0x13029 ADD %R12,%R10 |
(207) 0x1302c MOV %R10,%R13 |
(207) 0x1302f NOP |
(208) 0x13030 MOV 0x80(%RBP),%RAX |
(208) 0x13037 CMP %R9,(%RAX,%R13,8) |
(208) 0x1303b JE 13100 |
(208) 0x13041 INC %R13 |
(208) 0x13044 CMP %R13,%RCX |
(208) 0x13047 JNE 13030 |
(207) 0x13049 NOPL (%RAX) |
(207) 0x13050 MOV -0x58(%RBP),%RDX |
(207) 0x13054 CMP -0x88(%RBP),%RDX |
(207) 0x1305b JGE 13740 |
(207) 0x13061 MOV 0x80(%RBP),%RAX |
(207) 0x13068 MOV %R9,(%RAX,%RDX,8) |
(207) 0x1306c MOV -0x30(%RBP),%R11 |
(207) 0x13070 MOV 0x10(%RBP),%RAX |
(207) 0x13074 VMOVQ (%RAX,%R11,8),%XMM0 |
(207) 0x1307a VMOVQ %XMM0,(%RSI,%RDX,8) |
(207) 0x1307f INC %RDX |
(207) 0x13082 MOV %RDX,-0x58(%RBP) |
(207) 0x13086 JMP 12dc0 |
(207) 0x130c0 MOV -0x30(%RBP),%R11 |
(207) 0x130c4 JMP 13171 |
(207) 0x13100 MOV -0x30(%RBP),%R11 |
(207) 0x13104 JMP 131f1 |
(207) 0x13140 KSHIFTLB $0x4,%K0,%K0 |
(207) 0x13146 KORB %K0,%K1,%K0 |
(207) 0x1314a KSHIFTLB $0x4,%K2,%K1 |
(207) 0x13150 KORB %K1,%K3,%K1 |
(207) 0x13154 KUNPCKBW %K0,%K1,%K0 |
(207) 0x13158 KMOVD %K0,%EAX |
(207) 0x1315c TZCNT %EAX,%EAX |
(207) 0x13160 ADD %RDX,%R13 |
(207) 0x13163 ADD %RAX,%R13 |
(207) 0x13166 MOV -0x30(%RBP),%R11 |
(207) 0x1316a MOV 0x88(%RBP),%RSI |
(207) 0x13171 MOV -0x48(%RBP),%R10 |
(207) 0x13175 MOV %RCX,%R8 |
(207) 0x13178 MOV 0x10(%RBP),%RAX |
(207) 0x1317c VMOVQ (%RAX,%R11,8),%XMM0 |
(207) 0x13182 MOV 0x70(%RBP),%RAX |
(207) 0x13186 VMOVQ %XMM0,(%RAX,%R13,8) |
(207) 0x1318c JMP 12dc7 |
(207) 0x131c0 KSHIFTLB $0x4,%K0,%K0 |
(207) 0x131c6 KORB %K0,%K1,%K0 |
(207) 0x131ca KSHIFTLB $0x4,%K2,%K1 |
(207) 0x131d0 KORB %K1,%K3,%K1 |
(207) 0x131d4 KUNPCKBW %K0,%K1,%K0 |
(207) 0x131d8 KMOVD %K0,%EAX |
(207) 0x131dc TZCNT %EAX,%EAX |
(207) 0x131e0 ADD %RDX,%R13 |
(207) 0x131e3 ADD %RAX,%R13 |
(207) 0x131e6 MOV -0x30(%RBP),%R11 |
(207) 0x131ea MOV 0x88(%RBP),%RSI |
(207) 0x131f1 MOV -0x48(%RBP),%R10 |
(207) 0x131f5 MOV %RCX,%R8 |
(207) 0x131f8 MOV 0x10(%RBP),%RAX |
(207) 0x131fc VMOVQ (%RAX,%R11,8),%XMM0 |
(207) 0x13202 VMOVQ %XMM0,(%RSI,%R13,8) |
(207) 0x13208 JMP 12dc7 |
0x13240 MOVQ $0,-0x88(%RBP) |
0x1324b TEST %R14,%R14 |
0x1324e JLE 12a52 |
0x13254 LEA -0x1(%R14),%RAX |
0x13258 MOV %RAX,-0x90(%RBP) |
0x1325f MOV %R13D,%EAX |
0x13262 AND $0x7f,%EAX |
0x13265 MOV $0x80,%ECX |
0x1326a SUB %EAX,%ECX |
0x1326c SHR $0x3,%ECX |
0x1326f CMP %RCX,%R12 |
0x13272 CMOVB %R12,%RCX |
0x13276 LEA (%R13,%RCX,8),%RDX |
0x1327b MOV %R12,%RAX |
0x1327e SUB %RCX,%RAX |
0x13281 AND $-0x10,%RAX |
0x13285 ADD %RCX,%RAX |
0x13288 MOV %RAX,-0xa0(%RBP) |
0x1328f MOV %R12,%R9 |
0x13292 MOVQ $0,-0x58(%RBP) |
0x1329a XOR %EDI,%EDI |
0x1329c JMP 132f1 |
(217) 0x132c0 MOV -0x88(%RBP),%RSI |
(217) 0x132c7 MOV -0x58(%RBP),%R8 |
(217) 0x132cb MOV %RAX,(%RSI,%R8,8) |
(217) 0x132cf MOV -0x78(%RBP),%RAX |
(217) 0x132d3 VMOVQ %XMM0,(%RAX,%R8,8) |
(217) 0x132d9 INC %R8 |
(217) 0x132dc MOV %R8,-0x58(%RBP) |
(217) 0x132e0 CMP -0x90(%RBP),%RDI |
(217) 0x132e7 LEA 0x1(%RDI),%RDI |
(217) 0x132eb JE 134c0 |
(217) 0x132f1 MOV %R9,-0x60(%RBP) |
(217) 0x132f5 LEA (%R11,%RDI,1),%R8 |
(217) 0x132f9 TEST %R12,%R12 |
(217) 0x132fc JLE 133e0 |
(217) 0x13302 MOV (%RBX,%R8,8),%R9 |
(217) 0x13306 MOV %R13D,%EAX |
(217) 0x13309 AND $0x7f,%EAX |
(217) 0x1330c MOV $0x80,%ESI |
(217) 0x13311 SUB %EAX,%ESI |
(217) 0x13313 SHR $0x3,%ESI |
(217) 0x13316 CMP %RSI,%R12 |
(217) 0x13319 MOV %RSI,%RAX |
(217) 0x1331c CMOVB %R12,%RAX |
(217) 0x13320 TEST %RAX,%RAX |
(217) 0x13323 JE 13343 |
(217) 0x13325 XOR %R10D,%R10D |
(217) 0x13328 NOPL (%RAX,%RAX,1) |
(220) 0x13330 CMP %R9,(%R13,%R10,8) |
(220) 0x13335 JE 1346a |
(220) 0x1333b INC %R10 |
(220) 0x1333e CMP %R10,%RCX |
(220) 0x13341 JNE 13330 |
(217) 0x13343 CMP %RSI,%R12 |
(217) 0x13346 JBE 133e0 |
(217) 0x1334c MOV %R12,%R11 |
(217) 0x1334f SUB %RAX,%R11 |
(217) 0x13352 MOV %R11,%RSI |
(217) 0x13355 AND $-0x10,%RSI |
(217) 0x13359 JE 133aa |
(217) 0x1335b LEA -0x1(%RSI),%RAX |
(217) 0x1335f VPBROADCASTQ %R9,%YMM0 |
(217) 0x13365 XOR %R10D,%R10D |
(217) 0x13368 NOPL (%RAX,%RAX,1) |
(219) 0x13370 VPCMPEQQ 0x20(%RDX,%R10,8),%YMM0,%K0 |
(219) 0x13378 VPCMPEQQ (%RDX,%R10,8),%YMM0,%K1 |
(219) 0x1337f VPCMPEQQ 0x60(%RDX,%R10,8),%YMM0,%K2 |
(219) 0x13387 VPCMPEQQ 0x40(%RDX,%R10,8),%YMM0,%K3 |
(219) 0x1338f KORB %K2,%K3,%K4 |
(219) 0x13393 KORB %K0,%K1,%K5 |
(219) 0x13397 KORTESTB %K4,%K5 |
(219) 0x1339b JNE 13440 |
(219) 0x133a1 ADD $0x10,%R10 |
(219) 0x133a5 CMP %RAX,%R10 |
(219) 0x133a8 JBE 13370 |
(217) 0x133aa CMP %R11,%RSI |
(217) 0x133ad MOV -0x30(%RBP),%R11 |
(217) 0x133b1 JAE 133e0 |
(217) 0x133b3 MOV -0xa0(%RBP),%R10 |
(217) 0x133ba NOPW (%RAX,%RAX,1) |
(218) 0x133c0 CMP %R9,(%R13,%R10,8) |
(218) 0x133c5 JE 1346a |
(218) 0x133cb INC %R10 |
(218) 0x133ce CMP %R10,%R12 |
(218) 0x133d1 JNE 133c0 |
(217) 0x133d3 NOPW %CS:(%RAX,%RAX,1) |
(217) 0x133e0 MOV (%RBX,%R8,8),%RAX |
(217) 0x133e4 MOV 0x10(%RBP),%RSI |
(217) 0x133e8 VMOVQ (%RSI,%R8,8),%XMM0 |
(217) 0x133ee MOV -0x60(%RBP),%R9 |
(217) 0x133f2 CMP -0x80(%RBP),%R9 |
(217) 0x133f6 JGE 132c0 |
(217) 0x133fc MOV %RAX,(%R13,%R9,8) |
(217) 0x13401 MOV -0x98(%RBP),%RAX |
(217) 0x13408 VMOVQ %XMM0,(%RAX,%R9,8) |
(217) 0x1340e INC %R9 |
(217) 0x13411 JMP 132e0 |
(217) 0x13440 KSHIFTLB $0x4,%K0,%K0 |
(217) 0x13446 KORB %K0,%K1,%K0 |
(217) 0x1344a KSHIFTLB $0x4,%K2,%K1 |
(217) 0x13450 KORB %K1,%K3,%K1 |
(217) 0x13454 KUNPCKBW %K0,%K1,%K0 |
(217) 0x13458 KMOVD %K0,%EAX |
(217) 0x1345c TZCNT %EAX,%EAX |
(217) 0x13460 ADD %RCX,%R10 |
(217) 0x13463 ADD %RAX,%R10 |
(217) 0x13466 MOV -0x30(%RBP),%R11 |
(217) 0x1346a MOV 0x10(%RBP),%RAX |
(217) 0x1346e VMOVQ (%RAX,%R8,8),%XMM0 |
(217) 0x13474 MOV -0x98(%RBP),%RAX |
(217) 0x1347b VMOVQ %XMM0,(%RAX,%R10,8) |
(217) 0x13481 MOV -0x60(%RBP),%R9 |
(217) 0x13485 JMP 132e0 |
0x134c0 ADD %R14,%R11 |
0x134c3 MOV -0x58(%RBP),%RCX |
0x134c7 LEA (%R9,%RCX,1),%R12 |
0x134cb MOV 0x48(%RBP),%RAX |
0x134cf MOV -0x48(%RBP),%R13 |
0x134d3 MOV %R12,(%RAX,%R13,8) |
0x134d7 MOV %RCX,-0x58(%RBP) |
0x134db TEST %RCX,%RCX |
0x134de JE 13684 |
0x134e4 MOV %R9,-0x60(%RBP) |
0x134e8 MOV %R11,-0x30(%RBP) |
0x134ec MOV 0x38(%RBP),%RAX |
0x134f0 MOV (%RAX,%R13,8),%RDI |
0x134f4 LEA (,%R12,8),%R14 |
0x134fc MOV %R14,%RSI |
0x134ff VZEROUPPER |
0x13502 CALL 30e0 <hypre_ReAlloc@plt> |
0x13507 MOV 0x38(%RBP),%RCX |
0x1350b MOV %RAX,(%RCX,%R13,8) |
0x1350f MOV 0x40(%RBP),%RAX |
0x13513 MOV (%RAX,%R13,8),%RDI |
0x13517 MOV %R14,%RSI |
0x1351a CALL 30e0 <hypre_ReAlloc@plt> |
0x1351f MOV 0x40(%RBP),%RCX |
0x13523 MOV %RAX,(%RCX,%R13,8) |
0x13527 MOV 0x50(%RBP),%RCX |
0x1352b MOV %R12,(%RCX,%R13,8) |
0x1352f MOV -0x58(%RBP),%R11 |
0x13533 TEST %R11,%R11 |
0x13536 JLE 13680 |
0x1353c MOV 0x38(%RBP),%RCX |
0x13540 MOV (%RCX,%R13,8),%RCX |
0x13544 MOV -0x88(%RBP),%R12 |
0x1354b LEA -0x8(%R12,%R11,8),%RDX |
0x13550 MOV -0x60(%RBP),%R10 |
0x13554 LEA (%RCX,%R10,8),%RDI |
0x13558 CMP %RDI,%RDX |
0x1355b SETAE %DL |
0x1355e LEA -0x1(%R11,%R10,1),%RSI |
0x13563 LEA (%RCX,%RSI,8),%RCX |
0x13567 CMP %R12,%RCX |
0x1356a SETAE %R8B |
0x1356e MOV -0x78(%RBP),%R9 |
0x13572 LEA -0x8(%R9,%R11,8),%RCX |
0x13577 LEA (%RAX,%R10,8),%R14 |
0x1357b CMP %R14,%RCX |
0x1357e SETB %CL |
0x13581 LEA (%RAX,%RSI,8),%RAX |
0x13585 CMP %R9,%RAX |
0x13588 SETB %AL |
0x1358b TEST %R8B,%DL |
0x1358e MOV %R11,%R13 |
0x13591 JNE 136c0 |
0x13597 OR %AL,%CL |
0x13599 JE 136c0 |
0x1359f CMP $0xd,%R13 |
0x135a3 JB 137c0 |
0x135a9 SAL $0x3,%R13 |
0x135ad MOV %R12,%RSI |
0x135b0 MOV %R13,%RDX |
0x135b3 CALL 3350 <__intel_avx_rep_memcpy@plt> |
0x135b8 MOV %R14,%RDI |
0x135bb MOV -0x78(%RBP),%RSI |
0x135bf MOV %R13,%RDX |
0x135c2 CALL 3350 <__intel_avx_rep_memcpy@plt> |
0x135c7 MOV -0x30(%RBP),%R11 |
0x135cb JMP 1386b |
0x13600 MOV %RCX,%R14 |
0x13603 MOV %R8,%R9 |
0x13606 MOV -0x40(%RBP),%RCX |
0x1360a JMP 137a2 |
0x13640 MOV -0xa0(%RBP),%R11 |
0x13647 MOV -0x68(%RBP),%RDI |
0x1364b MOV -0x50(%RBP),%RSI |
0x1364f MOV -0x70(%RBP),%RDX |
0x13653 MOV -0x40(%RBP),%RCX |
0x13657 JMP 1379a |
0x13680 MOV -0x30(%RBP),%R11 |
0x13684 MOV -0x88(%RBP),%R12 |
0x1368b TEST %R12,%R12 |
0x1368e JNE 1386b |
0x13694 JMP 1388d |
0x136c0 XOR %EAX,%EAX |
0x136c2 MOV -0x30(%RBP),%R11 |
0x136c6 MOV -0x78(%RBP),%RDX |
0x136ca NOPW (%RAX,%RAX,1) |
(214) 0x136d0 MOV (%R12,%RAX,8),%RCX |
(214) 0x136d4 MOV %RCX,(%RDI,%RAX,8) |
(214) 0x136d8 VMOVQ (%RDX,%RAX,8),%XMM0 |
(214) 0x136dd VMOVQ %XMM0,(%R14,%RAX,8) |
(214) 0x136e3 INC %RAX |
(214) 0x136e6 CMP %RAX,%R13 |
(214) 0x136e9 JNE 136d0 |
0x136eb JMP 1386b |
0x13700 MOV $0xd70,%ESI |
0x13705 MOV $0x1,%EDX |
0x1370a LEA 0x5b62(%RIP),%RDI |
0x13711 XOR %ECX,%ECX |
0x13713 VZEROUPPER |
0x13716 CALL 3420 <hypre_error_handler@plt> |
0x1371b MOV 0xd0(%RBP),%RAX |
0x13722 LOCK INCQ (%RAX) |
0x13726 CMPQ $0,0xc0(%RBP) |
0x1372e JE 13782 |
0x13730 LEA 0x5cba(%RIP),%RDI |
0x13737 JMP 13777 |
0x13740 MOV $0xd4e,%ESI |
0x13745 MOV $0x1,%EDX |
0x1374a LEA 0x5b22(%RIP),%RDI |
0x13751 XOR %ECX,%ECX |
0x13753 VZEROUPPER |
0x13756 CALL 3420 <hypre_error_handler@plt> |
0x1375b MOV 0xd0(%RBP),%RAX |
0x13762 LOCK INCQ (%RAX) |
0x13766 CMPQ $0,0xc0(%RBP) |
0x1376e JE 13782 |
0x13770 LEA 0x5c54(%RIP),%RDI |
0x13777 MOV -0x60(%RBP),%RSI |
0x1377b XOR %EAX,%EAX |
0x1377d CALL 3410 <hypre_printf@plt> |
0x13782 MOV -0x68(%RBP),%RDI |
0x13786 MOV -0x50(%RBP),%RSI |
0x1378a MOV -0x70(%RBP),%RDX |
0x1378e MOV -0x40(%RBP),%RCX |
0x13792 MOV -0x30(%RBP),%R11 |
0x13796 MOV -0x48(%RBP),%R10 |
0x1379a MOV -0x58(%RBP),%R9 |
0x1379e MOV -0x80(%RBP),%R14 |
0x137a2 MOV 0x18(%RBP),%R8 |
0x137a6 MOV 0x38(%R8),%RAX |
0x137aa MOV %R14,(%RAX,%R10,8) |
0x137ae MOV 0x40(%R8),%RAX |
0x137b2 MOV %R9,(%RAX,%R10,8) |
0x137b6 JMP 12994 |
0x137c0 MOV %R13,%RAX |
0x137c3 AND $-0x4,%RAX |
0x137c7 MOV -0x30(%RBP),%R11 |
0x137cb JE 13840 |
0x137cd LEA -0x1(%RAX),%RCX |
0x137d1 XOR %EDX,%EDX |
0x137d3 MOV -0x78(%RBP),%RSI |
0x137d7 NOPW (%RAX,%RAX,1) |
(216) 0x137e0 VMOVUPS (%R12,%RDX,8),%YMM0 |
(216) 0x137e6 VMOVUPS %YMM0,(%RDI,%RDX,8) |
(216) 0x137eb VMOVDQU (%RSI,%RDX,8),%YMM0 |
(216) 0x137f0 VMOVDQU %YMM0,(%R14,%RDX,8) |
(216) 0x137f6 ADD $0x4,%RDX |
(216) 0x137fa CMP %RCX,%RDX |
(216) 0x137fd JLE 137e0 |
0x137ff CMP %RAX,%R13 |
0x13802 JNE 13842 |
0x13804 JMP 1386b |
0x13840 XOR %EAX,%EAX |
0x13842 MOV -0x78(%RBP),%RDX |
0x13846 NOPW %CS:(%RAX,%RAX,1) |
(215) 0x13850 MOV (%R12,%RAX,8),%RCX |
(215) 0x13854 MOV %RCX,(%RDI,%RAX,8) |
(215) 0x13858 VMOVQ (%RDX,%RAX,8),%XMM0 |
(215) 0x1385d VMOVQ %XMM0,(%R14,%RAX,8) |
(215) 0x13863 INC %RAX |
(215) 0x13866 CMP %RAX,%R13 |
(215) 0x13869 JNE 13850 |
0x1386b MOV %R12,%RDI |
0x1386e MOV %R11,%R14 |
0x13871 VZEROUPPER |
0x13874 CALL 31c0 <hypre_Free@plt> |
0x13879 MOV -0x78(%RBP),%RDI |
0x1387d CALL 31c0 <hypre_Free@plt> |
0x13882 MOV %R14,%R11 |
0x13885 MOVQ $0,-0x78(%RBP) |
0x1388d MOV -0x68(%RBP),%RDI |
0x13891 MOV -0x50(%RBP),%RSI |
0x13895 MOV -0x70(%RBP),%RDX |
0x13899 MOV -0x40(%RBP),%RCX |
0x1389d JMP 12994 |
/home/eoseret/qaas_runs_CPU_9468/171-716-5699/intel/AMG/build/AMG/AMG/IJ_mv/IJMatrix_parcsr.c: 3262 - 3484 |
-------------------------------------------------------------------------------- |
3262: if (my_thread_num < rest) |
[...] |
3291: for (ii=ns; ii < ne; ii++) |
3292: { |
3293: row = rows[ii]; |
3294: n = ncols[ii]; |
3295: /* processor owns the row */ |
3296: if (row >= row_partitioning[pstart] && row < row_partitioning[pstart+1]) |
3297: { |
3298: row_local = row - row_partitioning[pstart]; |
3299: /* compute local row number */ |
3300: if (need_aux) |
3301: { |
3302: local_j = aux_j[row_local]; |
3303: local_data = aux_data[row_local]; |
3304: space = row_space[row_local]; |
3305: old_size = row_length[row_local]; |
3306: size = space - old_size; |
3307: if (size < n) |
3308: { |
3309: size = n - size; |
3310: tmp_j = hypre_CTAlloc(HYPRE_Int,size); |
3311: tmp_data = hypre_CTAlloc(HYPRE_Complex,size); |
3312: } |
3313: tmp_indx = 0; |
3314: not_found = 1; |
3315: size = old_size; |
3316: for (i=0; i < n; i++) |
3317: { |
3318: for (j=0; j < old_size; j++) |
3319: { |
3320: if (local_j[j] == cols[indx]) |
3321: { |
3322: local_data[j] = values[indx]; |
[...] |
3329: if (size < space) |
3330: { |
3331: local_j[size] = cols[indx]; |
3332: local_data[size++] = values[indx]; |
3333: } |
3334: else |
3335: { |
3336: tmp_j[tmp_indx] = cols[indx]; |
3337: tmp_data[tmp_indx++] = values[indx]; |
[...] |
3344: row_length[row_local] = size+tmp_indx; |
3345: |
3346: if (tmp_indx) |
3347: { |
3348: aux_j[row_local] = hypre_TReAlloc(aux_j[row_local],HYPRE_Int, |
3349: size+tmp_indx); |
3350: aux_data[row_local] = hypre_TReAlloc(aux_data[row_local], |
3351: HYPRE_Complex,size+tmp_indx); |
3352: row_space[row_local] = size+tmp_indx; |
3353: local_j = aux_j[row_local]; |
[...] |
3359: for (i=0; i < tmp_indx; i++) |
3360: { |
3361: local_j[cnt] = tmp_j[i]; |
3362: local_data[cnt++] = tmp_data[i]; |
3363: } |
3364: |
3365: if (tmp_j) |
3366: { |
3367: hypre_TFree(tmp_j); |
3368: hypre_TFree(tmp_data); |
[...] |
3376: offd_indx = hypre_AuxParCSRMatrixIndxOffd(aux_matrix)[row_local]; |
3377: diag_indx = hypre_AuxParCSRMatrixIndxDiag(aux_matrix)[row_local]; |
3378: cnt_diag = diag_indx; |
3379: cnt_offd = offd_indx; |
3380: diag_space = diag_i[row_local+1]; |
3381: offd_space = offd_i[row_local+1]; |
3382: not_found = 1; |
3383: for (i=0; i < n; i++) |
3384: { |
3385: if (cols[indx] < col_0 || cols[indx] > col_n) |
3386: /* insert into offd */ |
3387: { |
3388: for (j=offd_i[row_local]; j < offd_indx; j++) |
3389: { |
3390: if (offd_j[j] == cols[indx]) |
3391: { |
3392: offd_data[j] = values[indx]; |
[...] |
3399: if (cnt_offd < offd_space) |
3400: { |
3401: offd_j[cnt_offd] = cols[indx]; |
3402: offd_data[cnt_offd++] = values[indx]; |
3403: } |
3404: else |
3405: { |
3406: hypre_error(HYPRE_ERROR_GENERIC); |
3407: #ifdef HYPRE_USING_OPENMP |
3408: #pragma omp atomic |
3409: #endif |
3410: error_flag++; |
3411: if (print_level) |
3412: hypre_printf("Error in row %d ! Too many elements!\n", |
[...] |
3422: for (j=diag_i[row_local]; j < diag_indx; j++) |
3423: { |
3424: if (diag_j[j] == cols[indx]) |
3425: { |
3426: diag_data[j] = values[indx]; |
[...] |
3433: if (cnt_diag < diag_space) |
3434: { |
3435: diag_j[cnt_diag] = cols[indx]; |
3436: diag_data[cnt_diag++] = values[indx]; |
3437: } |
3438: else |
3439: { |
3440: hypre_error(HYPRE_ERROR_GENERIC); |
3441: #ifdef HYPRE_USING_OPENMP |
3442: #pragma omp atomic |
3443: #endif |
3444: error_flag++; |
3445: if (print_level) |
3446: hypre_printf("Error in row %d ! Too many elements !\n", |
[...] |
3454: indx++; |
3455: } |
3456: |
3457: hypre_AuxParCSRMatrixIndxDiag(aux_matrix)[row_local] = cnt_diag; |
3458: hypre_AuxParCSRMatrixIndxOffd(aux_matrix)[row_local] = cnt_offd; |
[...] |
3466: indx += n; |
3467: if (aux_matrix) |
3468: { |
3469: col_indx = 0; |
3470: for (i=0; i < off_proc_i_indx; i=i+2) |
3471: { |
3472: row_len = off_proc_i[i+1]; |
3473: if (off_proc_i[i] == row) |
3474: { |
3475: for (j=0; j < n; j++) |
3476: { |
3477: cnt1 = col_indx; |
3478: for (k=0; k < row_len; k++) |
3479: { |
3480: if (off_proc_j[cnt1] == cols[j]) |
3481: { |
3482: off_proc_j[cnt1++] = -1; |
3483: /*cancel_indx++;*/ |
3484: offproc_cnt[my_thread_num]++; |
Path / |
Metric | Value |
---|---|
CQA speedup if no scalar integer | 1.00 |
CQA speedup if FP arith vectorized | 1.00 |
CQA speedup if fully vectorized | 21.33 |
CQA speedup if no inter-iteration dependency | NA |
CQA speedup if next bottleneck killed | 1.59 |
Bottlenecks | micro-operation queue, |
Function | hypre_IJMatrixSetValuesOMPParCSR.extracted.28 |
Source | IJMatrix_parcsr.c:3262-3262,IJMatrix_parcsr.c:3291-3296,IJMatrix_parcsr.c:3300-3307,IJMatrix_parcsr.c:3310-3311,IJMatrix_parcsr.c:3316-3316,IJMatrix_parcsr.c:3344-3353,IJMatrix_parcsr.c:3359-3362,IJMatrix_parcsr.c:3365-3368,IJMatrix_parcsr.c:3376-3377,IJMatrix_parcsr.c:3380-3383,IJMatrix_parcsr.c:3390-3392,IJMatrix_parcsr.c:3406-3406,IJMatrix_parcsr.c:3410-3412,IJMatrix_parcsr.c:3440-3440,IJMatrix_parcsr.c:3444-3446,IJMatrix_parcsr.c:3457-3458,IJMatrix_parcsr.c:3466-3467,IJMatrix_parcsr.c:3473-3475,IJMatrix_parcsr.c:3484-3484 |
Source loop unroll info | NA |
Source loop unroll confidence level | NA |
Unroll/vectorization loop type | NA |
Unroll factor | NA |
CQA cycles | 47.83 |
CQA cycles if no scalar integer | 47.83 |
CQA cycles if FP arith vectorized | 47.83 |
CQA cycles if fully vectorized | 2.24 |
Front-end cycles | 47.83 |
DIV/SQRT cycles | 14.60 |
P0 cycles | 14.60 |
P1 cycles | 30.00 |
P2 cycles | 30.00 |
P3 cycles | 23.50 |
P4 cycles | 14.60 |
P5 cycles | 14.60 |
P6 cycles | 23.50 |
P7 cycles | 23.50 |
P8 cycles | 23.50 |
P9 cycles | 14.60 |
P10 cycles | 30.00 |
P11 cycles | 0.00 |
Inter-iter dependencies cycles | NA |
FE+BE cycles (UFS) | 45.34 |
Stall cycles (UFS) | 0.00 |
Nb insns | 267.00 |
Nb uops | 287.00 |
Nb loads | 90.00 |
Nb stores | 36.00 |
Nb stack references | 28.00 |
FLOP/cycle | 0.00 |
Nb FLOP add-sub | 0.00 |
Nb FLOP mul | 0.00 |
Nb FLOP fma | 0.00 |
Nb FLOP div | 0.00 |
Nb FLOP rcp | 0.00 |
Nb FLOP sqrt | 0.00 |
Nb FLOP rsqrt | 0.00 |
Bytes/cycle | 20.93 |
Bytes prefetched | 0.00 |
Bytes loaded | 713.00 |
Bytes stored | 288.00 |
Stride 0 | NA |
Stride 1 | NA |
Stride n | NA |
Stride unknown | NA |
Stride indirect | NA |
Vectorization ratio all | 8.33 |
Vectorization ratio load | 0.00 |
Vectorization ratio store | 0.00 |
Vectorization ratio mul | NA |
Vectorization ratio add_sub | NA |
Vectorization ratio fma | NA |
Vectorization ratio div_sqrt | NA |
Vectorization ratio other | 20.83 |
Vector-efficiency ratio all | 13.05 |
Vector-efficiency ratio load | 10.94 |
Vector-efficiency ratio store | 11.98 |
Vector-efficiency ratio mul | NA |
Vector-efficiency ratio add_sub | NA |
Vector-efficiency ratio fma | NA |
Vector-efficiency ratio div_sqrt | NA |
Vector-efficiency ratio other | 14.65 |
Metric | Value |
---|---|
CQA speedup if no scalar integer | 1.00 |
CQA speedup if FP arith vectorized | 1.00 |
CQA speedup if fully vectorized | 21.33 |
CQA speedup if no inter-iteration dependency | NA |
CQA speedup if next bottleneck killed | 1.59 |
Bottlenecks | micro-operation queue, |
Function | hypre_IJMatrixSetValuesOMPParCSR.extracted.28 |
Source | IJMatrix_parcsr.c:3262-3262,IJMatrix_parcsr.c:3291-3296,IJMatrix_parcsr.c:3300-3307,IJMatrix_parcsr.c:3310-3311,IJMatrix_parcsr.c:3316-3316,IJMatrix_parcsr.c:3344-3353,IJMatrix_parcsr.c:3359-3362,IJMatrix_parcsr.c:3365-3368,IJMatrix_parcsr.c:3376-3377,IJMatrix_parcsr.c:3380-3383,IJMatrix_parcsr.c:3390-3392,IJMatrix_parcsr.c:3406-3406,IJMatrix_parcsr.c:3410-3412,IJMatrix_parcsr.c:3440-3440,IJMatrix_parcsr.c:3444-3446,IJMatrix_parcsr.c:3457-3458,IJMatrix_parcsr.c:3466-3467,IJMatrix_parcsr.c:3473-3475,IJMatrix_parcsr.c:3484-3484 |
Source loop unroll info | NA |
Source loop unroll confidence level | NA |
Unroll/vectorization loop type | NA |
Unroll factor | NA |
CQA cycles | 47.83 |
CQA cycles if no scalar integer | 47.83 |
CQA cycles if FP arith vectorized | 47.83 |
CQA cycles if fully vectorized | 2.24 |
Front-end cycles | 47.83 |
DIV/SQRT cycles | 14.60 |
P0 cycles | 14.60 |
P1 cycles | 30.00 |
P2 cycles | 30.00 |
P3 cycles | 23.50 |
P4 cycles | 14.60 |
P5 cycles | 14.60 |
P6 cycles | 23.50 |
P7 cycles | 23.50 |
P8 cycles | 23.50 |
P9 cycles | 14.60 |
P10 cycles | 30.00 |
P11 cycles | 0.00 |
Inter-iter dependencies cycles | NA |
FE+BE cycles (UFS) | 45.34 |
Stall cycles (UFS) | 0.00 |
Nb insns | 267.00 |
Nb uops | 287.00 |
Nb loads | 90.00 |
Nb stores | 36.00 |
Nb stack references | 28.00 |
FLOP/cycle | 0.00 |
Nb FLOP add-sub | 0.00 |
Nb FLOP mul | 0.00 |
Nb FLOP fma | 0.00 |
Nb FLOP div | 0.00 |
Nb FLOP rcp | 0.00 |
Nb FLOP sqrt | 0.00 |
Nb FLOP rsqrt | 0.00 |
Bytes/cycle | 20.93 |
Bytes prefetched | 0.00 |
Bytes loaded | 713.00 |
Bytes stored | 288.00 |
Stride 0 | NA |
Stride 1 | NA |
Stride n | NA |
Stride unknown | NA |
Stride indirect | NA |
Vectorization ratio all | 8.33 |
Vectorization ratio load | 0.00 |
Vectorization ratio store | 0.00 |
Vectorization ratio mul | NA |
Vectorization ratio add_sub | NA |
Vectorization ratio fma | NA |
Vectorization ratio div_sqrt | NA |
Vectorization ratio other | 20.83 |
Vector-efficiency ratio all | 13.05 |
Vector-efficiency ratio load | 10.94 |
Vector-efficiency ratio store | 11.98 |
Vector-efficiency ratio mul | NA |
Vector-efficiency ratio add_sub | NA |
Vector-efficiency ratio fma | NA |
Vector-efficiency ratio div_sqrt | NA |
Vector-efficiency ratio other | 14.65 |
Path / |
Function | hypre_IJMatrixSetValuesOMPParCSR.extracted.28 |
Source file and lines | IJMatrix_parcsr.c:3262-3484 |
Module | libIJ_mv.so |
nb instructions | 267 |
nb uops | 287 |
loop length | 1136 |
used x86 registers | 13 |
used mmx registers | 0 |
used xmm registers | 0 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 28 |
micro-operation queue | 47.83 cycles |
front end | 47.83 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 14.60 | 14.60 | 30.00 | 30.00 | 23.50 | 14.60 | 14.60 | 23.50 | 23.50 | 23.50 | 14.60 | 30.00 |
cycles | 14.60 | 14.60 | 30.00 | 30.00 | 23.50 | 14.60 | 14.60 | 23.50 | 23.50 | 23.50 | 14.60 | 30.00 |
Cycles executing div or sqrt instructions | NA |
FE+BE cycles | 45.34 |
Stall cycles | 0.00 |
Front-end | 47.83 |
Dispatch | 30.00 |
Overall L1 | 47.83 |
all | 8% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 20% |
all | 13% |
load | 10% |
store | 11% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 14% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
MOV -0x68(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x70(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
INC %RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RCX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 138c0 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1200> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV (%RDI,%RDX,8),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RSI,%RDX,8),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R10,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
SUB (%RAX),%R10 | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JL 12a80 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x3c0> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x60(%RBP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP 0x8(%RAX),%R8 | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JGE 12a80 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x3c0> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMPQ $0,0x58(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
MOV %R10,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JE 12d40 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x680> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x40(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x50(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %R12,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R14,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %RAX,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RDX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R11,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JLE 13240 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xb80> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV $0x8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RDI,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 3360 <hypre_CAlloc@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV $0x8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x60(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CALL 3360 <hypre_CAlloc@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JG 13254 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xb94> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x48(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R12,(%RAX,%RCX,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JMP 13684 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xfc4> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
ADD %R14,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMPB $0,-0x31(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JNE 12998 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x2d8> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 12998 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x2d8> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R11,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x98(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
DEC %RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SHR $0x1,%RAX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %RAX,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
DEC %R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 12ad1 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x411> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
MOV 0x18(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%RCX),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x40(%RCX),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RCX,%R10,8),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 13600 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xf40> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %RDX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x60(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x8(%RAX,%R10,8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x78(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x8(%RAX,%R10,8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
XOR %EDI,%EDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R8,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
LEA (%R11,%R14,1),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,-0xa0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x88(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R8,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
JMP 12dd6 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x716> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
MOVQ $0,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 12a52 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x392> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%R14),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R13D,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $0x7f,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV $0x80,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SUB %EAX,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SHR $0x3,%ECX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
CMP %RCX,%R12 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMOVB %R12,%RCX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
LEA (%R13,%RCX,8),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R12,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %RCX,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
AND $-0x10,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
ADD %RCX,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RAX,-0xa0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R12,%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOVQ $0,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
XOR %EDI,%EDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 132f1 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xc31> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
ADD %R14,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x58(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%R9,%RCX,1),%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x48(%RBP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R12,(%RAX,%R13,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JE 13684 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xfc4> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R9,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R11,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R13,8),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (,%R12,8),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R14,%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 30e0 <hypre_ReAlloc@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0x38(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,(%RCX,%R13,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x40(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R13,8),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R14,%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 30e0 <hypre_ReAlloc@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,(%RCX,%R13,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x50(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R12,(%RCX,%R13,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x58(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R11,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 13680 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xfc0> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x38(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RCX,%R13,8),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x88(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA -0x8(%R12,%R11,8),%RDX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV -0x60(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RCX,%R10,8),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %RDI,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SETAE %DL | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
LEA -0x1(%R11,%R10,1),%RSI | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA (%RCX,%RSI,8),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %R12,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SETAE %R8B | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV -0x78(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA -0x8(%R9,%R11,8),%RCX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA (%RAX,%R10,8),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %R14,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SETB %CL | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
LEA (%RAX,%RSI,8),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %R9,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SETB %AL | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
TEST %R8B,%DL | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
MOV %R11,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
JNE 136c0 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1000> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
OR %AL,%CL | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
JE 136c0 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1000> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP $0xd,%R13 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JB 137c0 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1100> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
SAL $0x3,%R13 | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %R12,%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R13,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 3350 <__intel_avx_rep_memcpy@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %R14,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x78(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R13,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 3350 <__intel_avx_rep_memcpy@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 1386b <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11ab> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV %RCX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R8,%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 137a2 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10e2> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV -0xa0(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x68(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x70(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 1379a <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10da> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x88(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R12,%R12 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JNE 1386b <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11ab> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 1388d <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11cd> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x78(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 1386b <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11ab> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV $0xd70,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA 0x5b62(%RIP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 3420 <hypre_error_handler@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0xd0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LOCK INCQ (%RAX) | 3 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
CMPQ $0,0xc0(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JE 13782 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10c2> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA 0x5cba(%RIP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 13777 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10b7> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
MOV $0xd4e,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA 0x5b22(%RIP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 3420 <hypre_error_handler@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0xd0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LOCK INCQ (%RAX) | 3 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
CMPQ $0,0xc0(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JE 13782 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10c2> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA 0x5c54(%RIP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x60(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 3410 <hypre_printf@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x68(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x70(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x48(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x58(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x80(%RBP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x18(%RBP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%R8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R14,(%RAX,%R10,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x40(%R8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R9,(%RAX,%R10,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JMP 12994 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x2d4> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV %R13,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $-0x4,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JE 13840 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1180> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%RAX),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x78(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %RAX,%R13 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 13842 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1182> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 1386b <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11ab> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x78(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R12,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R11,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 31c0 <hypre_Free@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x78(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CALL 31c0 <hypre_Free@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %R14,%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOVQ $0,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x68(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x70(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 12994 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x2d4> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
Function | hypre_IJMatrixSetValuesOMPParCSR.extracted.28 |
Source file and lines | IJMatrix_parcsr.c:3262-3484 |
Module | libIJ_mv.so |
nb instructions | 267 |
nb uops | 287 |
loop length | 1136 |
used x86 registers | 13 |
used mmx registers | 0 |
used xmm registers | 0 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 28 |
micro-operation queue | 47.83 cycles |
front end | 47.83 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 14.60 | 14.60 | 30.00 | 30.00 | 23.50 | 14.60 | 14.60 | 23.50 | 23.50 | 23.50 | 14.60 | 30.00 |
cycles | 14.60 | 14.60 | 30.00 | 30.00 | 23.50 | 14.60 | 14.60 | 23.50 | 23.50 | 23.50 | 14.60 | 30.00 |
Cycles executing div or sqrt instructions | NA |
FE+BE cycles | 45.34 |
Stall cycles | 0.00 |
Front-end | 47.83 |
Dispatch | 30.00 |
Overall L1 | 47.83 |
all | 8% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 20% |
all | 13% |
load | 10% |
store | 11% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 14% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
MOV -0x68(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x70(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
INC %RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RCX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 138c0 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1200> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV (%RDI,%RDX,8),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RSI,%RDX,8),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R10,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
SUB (%RAX),%R10 | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JL 12a80 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x3c0> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x60(%RBP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP 0x8(%RAX),%R8 | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JGE 12a80 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x3c0> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMPQ $0,0x58(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
MOV %R10,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JE 12d40 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x680> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x40(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x50(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %R12,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R14,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %RAX,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RDX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R11,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JLE 13240 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xb80> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV $0x8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RDI,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 3360 <hypre_CAlloc@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV $0x8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x60(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CALL 3360 <hypre_CAlloc@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JG 13254 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xb94> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x48(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R12,(%RAX,%RCX,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JMP 13684 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xfc4> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
ADD %R14,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMPB $0,-0x31(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JNE 12998 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x2d8> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 12998 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x2d8> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R11,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x98(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
DEC %RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SHR $0x1,%RAX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %RAX,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
DEC %R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 12ad1 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x411> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
MOV 0x18(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%RCX),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x40(%RCX),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RCX,%R10,8),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R10,8),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 13600 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xf40> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %RDX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x60(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x8(%RAX,%R10,8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x78(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x8(%RAX,%R10,8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
XOR %EDI,%EDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R8,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
LEA (%R11,%R14,1),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,-0xa0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x88(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R8,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
JMP 12dd6 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x716> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
MOVQ $0,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 12a52 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x392> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%R14),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R13D,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $0x7f,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV $0x80,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SUB %EAX,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SHR $0x3,%ECX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
CMP %RCX,%R12 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMOVB %R12,%RCX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
LEA (%R13,%RCX,8),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R12,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %RCX,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
AND $-0x10,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
ADD %RCX,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RAX,-0xa0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R12,%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOVQ $0,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
XOR %EDI,%EDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 132f1 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xc31> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
ADD %R14,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x58(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%R9,%RCX,1),%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x48(%RBP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R12,(%RAX,%R13,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JE 13684 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xfc4> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R9,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R11,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x38(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R13,8),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (,%R12,8),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R14,%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 30e0 <hypre_ReAlloc@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0x38(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,(%RCX,%R13,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x40(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX,%R13,8),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R14,%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 30e0 <hypre_ReAlloc@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,(%RCX,%R13,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x50(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R12,(%RCX,%R13,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x58(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R11,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 13680 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0xfc0> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x38(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RCX,%R13,8),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x88(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA -0x8(%R12,%R11,8),%RDX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV -0x60(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RCX,%R10,8),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %RDI,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SETAE %DL | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
LEA -0x1(%R11,%R10,1),%RSI | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA (%RCX,%RSI,8),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %R12,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SETAE %R8B | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV -0x78(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA -0x8(%R9,%R11,8),%RCX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA (%RAX,%R10,8),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %R14,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SETB %CL | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
LEA (%RAX,%RSI,8),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %R9,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SETB %AL | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
TEST %R8B,%DL | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
MOV %R11,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
JNE 136c0 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1000> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
OR %AL,%CL | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
JE 136c0 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1000> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP $0xd,%R13 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JB 137c0 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1100> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
SAL $0x3,%R13 | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %R12,%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R13,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 3350 <__intel_avx_rep_memcpy@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %R14,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x78(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R13,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 3350 <__intel_avx_rep_memcpy@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 1386b <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11ab> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV %RCX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R8,%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 137a2 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10e2> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV -0xa0(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x68(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x70(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 1379a <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10da> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x88(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R12,%R12 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JNE 1386b <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11ab> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 1388d <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11cd> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x78(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 1386b <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11ab> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV $0xd70,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA 0x5b62(%RIP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 3420 <hypre_error_handler@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0xd0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LOCK INCQ (%RAX) | 3 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
CMPQ $0,0xc0(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JE 13782 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10c2> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA 0x5cba(%RIP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 13777 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10b7> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
MOV $0xd4e,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA 0x5b22(%RIP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 3420 <hypre_error_handler@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0xd0(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LOCK INCQ (%RAX) | 3 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
CMPQ $0,0xc0(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JE 13782 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x10c2> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA 0x5c54(%RIP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x60(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 3410 <hypre_printf@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x68(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x70(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x48(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x58(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x80(%RBP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x18(%RBP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%R8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R14,(%RAX,%R10,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x40(%R8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R9,(%RAX,%R10,8) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JMP 12994 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x2d4> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV %R13,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $-0x4,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JE 13840 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1180> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%RAX),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x78(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP %RAX,%R13 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 13842 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x1182> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 1386b <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x11ab> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x78(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R12,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R11,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 31c0 <hypre_Free@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x78(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CALL 31c0 <hypre_Free@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %R14,%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOVQ $0,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x68(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x50(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x70(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 12994 <hypre_IJMatrixSetValuesOMPParCSR.extracted.28+0x2d4> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |