Function: _Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buf ... | Module: exec | Source: advec_mom.cpp:167-172 [...] | Coverage: 2.81% |
---|
Function: _Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buf ... | Module: exec | Source: advec_mom.cpp:167-172 [...] | Coverage: 2.81% |
---|
/beegfs/hackathon/users/eoseret/qaas_runs/170-854-8685/intel/CloverLeafCXX/build/CloverLeafCXX/src/omp/advec_mom.cpp: 167 - 172 |
-------------------------------------------------------------------------------- |
167: #pragma omp parallel for simd collapse(2) |
168: for (int j = (y_min - 1 + 1); j < (y_max + 2 + 2); j++) { |
169: for (int i = (x_min + 1); i < (x_max + 1 + 2); i++) { |
170: node_mass_post(i, j) = 0.25 * (density1(i + 0, j - 1) * post_vol(i + 0, j - 1) + density1(i, j) * post_vol(i, j) + |
171: density1(i - 1, j - 1) * post_vol(i - 1, j - 1) + density1(i - 1, j + 0) * post_vol(i - 1, j + 0)); |
172: node_mass_pre(i, j) = node_mass_post(i, j) - node_flux(i + 0, j - 1) + node_flux(i, j); |
/beegfs/hackathon/users/eoseret/qaas_runs/170-854-8685/intel/CloverLeafCXX/build/CloverLeafCXX/src/omp/context.h: 69 - 69 |
-------------------------------------------------------------------------------- |
69: T &operator()(size_t i, size_t j) const { return data[i + j * sizeX]; } |
0x4398f0 PUSH %RBP |
0x4398f1 MOV %RSP,%RBP |
0x4398f4 PUSH %R15 |
0x4398f6 PUSH %R14 |
0x4398f8 PUSH %R13 |
0x4398fa PUSH %R12 |
0x4398fc PUSH %RBX |
0x4398fd SUB $0x168,%RSP |
0x439904 MOV %RCX,%R15 |
0x439907 MOV 0x38(%RBP),%RAX |
0x43990b MOV 0x28(%RBP),%RBX |
0x43990f MOV 0x20(%RBP),%R10 |
0x439913 MOV 0x10(%RBP),%R14 |
0x439917 MOV 0x18(%RBP),%ECX |
0x43991a MOV %ECX,-0x2c(%RBP) |
0x43991d MOVL $0,-0x40(%RBP) |
0x439924 TEST %RAX,%RAX |
0x439927 JS 43a009 |
0x43992d MOV %R8,%R12 |
0x439930 MOV %RDX,%R13 |
0x439933 MOV %R9,-0x38(%RBP) |
0x439937 MOV (%RDI),%ESI |
0x439939 MOVQ $0,-0x60(%RBP) |
0x439941 MOV %RAX,-0x58(%RBP) |
0x439945 MOVQ $0x1,-0xb0(%RBP) |
0x439950 SUB $0x8,%RSP |
0x439954 LEA -0xb0(%RBP),%RAX |
0x43995b LEA -0x40(%RBP),%RCX |
0x43995f LEA -0x60(%RBP),%R8 |
0x439963 LEA -0x58(%RBP),%R9 |
0x439967 MOV $0x4aa900,%EDI |
0x43996c MOV %ESI,-0x3c(%RBP) |
0x43996f MOV $0x22,%EDX |
0x439974 PUSH $0x1 |
0x439976 PUSH $0x1 |
0x439978 PUSH %RAX |
0x439979 MOV %R10,-0x48(%RBP) |
0x43997d CALL 404240 <__kmpc_for_static_init_8@plt> |
0x439982 MOV -0x48(%RBP),%R11 |
0x439986 ADD $0x20,%RSP |
0x43998a MOV -0x60(%RBP),%RCX |
0x43998e MOV -0x58(%RBP),%RDX |
0x439992 CMP %RDX,%RCX |
0x439995 JA 439feb |
0x43999b MOV %RDX,%RAX |
0x43999e SUB %R11D,%EBX |
0x4399a1 MOV (%R13),%RDX |
0x4399a5 MOV 0x10(%R13),%R10 |
0x4399a9 MOV (%R14),%RSI |
0x4399ac MOV 0x10(%R14),%R13 |
0x4399b0 MOV (%R12),%RDI |
0x4399b4 MOV 0x10(%R12),%R8 |
0x4399b9 MOV (%R15),%R14 |
0x4399bc MOV 0x10(%R15),%R9 |
0x4399c0 MOV -0x38(%RBP),%R12 |
0x4399c4 MOV (%R12),%R15 |
0x4399c8 MOV 0x10(%R12),%R12 |
0x4399cd MOV %R12,-0x38(%RBP) |
0x4399d1 INC %RAX |
0x4399d4 MOV %RAX,-0xa8(%RBP) |
0x4399db SUB %RCX,%RAX |
0x4399de MOV $-0x2,%R12D |
0x4399e4 AND %RAX,%R12 |
0x4399e7 MOV %R9,-0x90(%RBP) |
0x4399ee MOV %R14,-0x50(%RBP) |
0x4399f2 MOV %RDI,-0xa0(%RBP) |
0x4399f9 MOV %R8,-0x98(%RBP) |
0x439a00 MOV %R15,-0x88(%RBP) |
0x439a07 JE 43a01b |
0x439a0d MOV %RAX,-0x68(%RBP) |
0x439a11 MOV %RBX,-0x80(%RBP) |
0x439a15 MOVD %EBX,%XMM0 |
0x439a19 PSHUFD $0x44,%XMM0,%XMM1 |
0x439a1e MOVD -0x2c(%RBP),%XMM0 |
0x439a23 PSHUFD $0x50,%XMM0,%XMM0 |
0x439a28 MOVDQA %XMM0,-0x180(%RBP) |
0x439a30 MOVD %R11D,%XMM0 |
0x439a35 PSHUFD $0x50,%XMM0,%XMM0 |
0x439a3a MOVDQA %XMM0,-0x170(%RBP) |
0x439a42 MOV %RDX,-0x78(%RBP) |
0x439a46 MOVQ %RDX,%XMM0 |
0x439a4b PSHUFD $0x44,%XMM0,%XMM0 |
0x439a50 MOVDQA %XMM0,-0x140(%RBP) |
0x439a58 MOVQ %R10,%XMM0 |
0x439a5d PSHUFD $0x44,%XMM0,%XMM0 |
0x439a62 MOVDQA %XMM0,-0x160(%RBP) |
0x439a6a MOV %RSI,-0x70(%RBP) |
0x439a6e MOVQ %RSI,%XMM0 |
0x439a73 PSHUFD $0x44,%XMM0,%XMM0 |
0x439a78 MOVDQA %XMM0,-0xd0(%RBP) |
0x439a80 MOVQ %R13,%XMM0 |
0x439a85 PSHUFD $0x44,%XMM0,%XMM0 |
0x439a8a MOVDQA %XMM0,-0x150(%RBP) |
0x439a92 MOVQ %RDI,%XMM0 |
0x439a97 PSHUFD $0x44,%XMM0,%XMM0 |
0x439a9c MOVDQA %XMM0,-0x130(%RBP) |
0x439aa4 MOVQ %R8,%XMM0 |
0x439aa9 PSHUFD $0x44,%XMM0,%XMM0 |
0x439aae MOVDQA %XMM0,-0x120(%RBP) |
0x439ab6 MOVQ %R14,%XMM0 |
0x439abb PSHUFD $0x44,%XMM0,%XMM0 |
0x439ac0 MOVDQA %XMM0,-0xc0(%RBP) |
0x439ac8 MOVQ %R9,%XMM0 |
0x439acd PSHUFD $0x44,%XMM0,%XMM0 |
0x439ad2 MOVDQA %XMM0,-0x110(%RBP) |
0x439ada MOVQ %R15,%XMM0 |
0x439adf PSHUFD $0x44,%XMM0,%XMM0 |
0x439ae4 MOVDQA %XMM0,-0x100(%RBP) |
0x439aec MOVQ -0x38(%RBP),%XMM0 |
0x439af1 PSHUFD $0x44,%XMM0,%XMM0 |
0x439af6 MOVDQA %XMM0,-0xf0(%RBP) |
0x439afe MOVQ %RCX,%XMM0 |
0x439b03 PSHUFD $0x44,%XMM0,%XMM13 |
0x439b09 PADDQ 0x51a7e(%RIP),%XMM13 |
0x439b12 XOR %R8D,%R8D |
0x439b15 MOVDQA %XMM1,-0x190(%RBP) |
0x439b1d PSHUFD $-0x12,%XMM1,%XMM0 |
0x439b22 MOVDQA %XMM0,-0xe0(%RBP) |
0x439b2a PXOR %XMM6,%XMM6 |
0x439b2e XCHG %AX,%AX |
(312) 0x439b30 MOVQ %XMM13,%RSI |
(312) 0x439b35 MOVDQA -0x190(%RBP),%XMM0 |
(312) 0x439b3d MOVQ %XMM0,%R9 |
(312) 0x439b42 MOV %RSI,%RAX |
(312) 0x439b45 XOR %EDX,%EDX |
(312) 0x439b47 DIV %R9 |
(312) 0x439b4a PSHUFD $-0x12,%XMM13,%XMM0 |
(312) 0x439b50 MOVQ %XMM0,%RDI |
(312) 0x439b55 MOVDQA -0xe0(%RBP),%XMM0 |
(312) 0x439b5d MOVQ %XMM0,%R14 |
(312) 0x439b62 MOVQ %RAX,%XMM0 |
(312) 0x439b67 MOV %RDI,%RAX |
(312) 0x439b6a XOR %EDX,%EDX |
(312) 0x439b6c DIV %R14 |
(312) 0x439b6f MOVQ %RAX,%XMM1 |
(312) 0x439b74 PUNPCKLQDQ %XMM1,%XMM0 |
(312) 0x439b78 PSHUFD $-0x18,%XMM0,%XMM15 |
(312) 0x439b7e PADDD -0x180(%RBP),%XMM15 |
(312) 0x439b87 MOV %RSI,%RAX |
(312) 0x439b8a CQTO |
(312) 0x439b8c IDIV %R9 |
(312) 0x439b8f MOVQ %RDX,%XMM0 |
(312) 0x439b94 MOV %RDI,%RAX |
(312) 0x439b97 CQTO |
(312) 0x439b99 IDIV %R14 |
(312) 0x439b9c MOVQ %RDX,%XMM1 |
(312) 0x439ba1 PXOR %XMM5,%XMM5 |
(312) 0x439ba5 PUNPCKLQDQ %XMM1,%XMM0 |
(312) 0x439ba9 MOVDQA %XMM15,%XMM2 |
(312) 0x439bae PCMPEQD %XMM1,%XMM1 |
(312) 0x439bb2 PADDD %XMM1,%XMM2 |
(312) 0x439bb6 PXOR %XMM12,%XMM12 |
(312) 0x439bbb PSHUFD $-0x18,%XMM0,%XMM0 |
(312) 0x439bc0 PCMPGTD %XMM2,%XMM12 |
(312) 0x439bc5 PUNPCKLDQ %XMM12,%XMM2 |
(312) 0x439bca MOVDQA -0x140(%RBP),%XMM7 |
(312) 0x439bd2 MOVDQA %XMM7,%XMM8 |
(312) 0x439bd7 PADDD -0x170(%RBP),%XMM0 |
(312) 0x439bdf PMULUDQ %XMM2,%XMM8 |
(312) 0x439be4 MOVDQA %XMM7,%XMM10 |
(312) 0x439be9 PSRLQ $0x20,%XMM10 |
(312) 0x439bef PCMPGTD %XMM0,%XMM5 |
(312) 0x439bf3 MOVDQA %XMM10,%XMM14 |
(312) 0x439bf8 PMULUDQ %XMM2,%XMM14 |
(312) 0x439bfd PUNPCKLDQ %XMM6,%XMM12 |
(312) 0x439c02 MOVDQA %XMM0,%XMM1 |
(312) 0x439c06 MOVDQA %XMM7,%XMM3 |
(312) 0x439c0a PMULUDQ %XMM12,%XMM3 |
(312) 0x439c0f MOVDQA -0xd0(%RBP),%XMM9 |
(312) 0x439c18 MOVDQA %XMM9,%XMM4 |
(312) 0x439c1d PUNPCKLDQ %XMM5,%XMM1 |
(312) 0x439c21 PSRLQ $0x20,%XMM4 |
(312) 0x439c26 MOVDQA %XMM4,%XMM11 |
(312) 0x439c2b PMULUDQ %XMM2,%XMM11 |
(312) 0x439c30 PADDQ %XMM14,%XMM3 |
(312) 0x439c35 MOVDQA %XMM9,%XMM5 |
(312) 0x439c3a PMULUDQ %XMM12,%XMM5 |
(312) 0x439c3f PADDQ %XMM11,%XMM5 |
(312) 0x439c44 PSLLQ $0x20,%XMM3 |
(312) 0x439c49 MOVDQA %XMM9,%XMM11 |
(312) 0x439c4e PMULUDQ %XMM2,%XMM11 |
(312) 0x439c53 PSLLQ $0x20,%XMM5 |
(312) 0x439c58 PADDQ %XMM8,%XMM3 |
(312) 0x439c5d PXOR %XMM8,%XMM8 |
(312) 0x439c62 PCMPGTD %XMM15,%XMM8 |
(312) 0x439c67 PUNPCKLDQ %XMM8,%XMM15 |
(312) 0x439c6c PADDQ %XMM11,%XMM5 |
(312) 0x439c71 PMULUDQ %XMM15,%XMM10 |
(312) 0x439c76 PUNPCKLDQ %XMM6,%XMM8 |
(312) 0x439c7b MOVDQA %XMM7,%XMM14 |
(312) 0x439c80 PMULUDQ %XMM8,%XMM14 |
(312) 0x439c85 PADDQ %XMM10,%XMM14 |
(312) 0x439c8a MOVDQA %XMM7,%XMM10 |
(312) 0x439c8f PMULUDQ %XMM15,%XMM10 |
(312) 0x439c94 PSLLQ $0x20,%XMM14 |
(312) 0x439c9a PADDQ %XMM10,%XMM14 |
(312) 0x439c9f PMULUDQ %XMM15,%XMM4 |
(312) 0x439ca4 MOVDQA %XMM9,%XMM10 |
(312) 0x439ca9 PMULUDQ %XMM8,%XMM10 |
(312) 0x439cae PADDQ %XMM4,%XMM10 |
(312) 0x439cb3 PMULUDQ %XMM15,%XMM9 |
(312) 0x439cb8 PSLLQ $0x20,%XMM10 |
(312) 0x439cbe PADDQ %XMM9,%XMM10 |
(312) 0x439cc3 MOVDQA %XMM14,%XMM4 |
(312) 0x439cc8 PADDQ %XMM1,%XMM4 |
(312) 0x439ccc PSLLQ $0x3,%XMM4 |
(312) 0x439cd1 MOVDQA -0x160(%RBP),%XMM11 |
(312) 0x439cda PADDQ %XMM11,%XMM4 |
(312) 0x439cdf MOVQ %XMM4,%RSI |
(312) 0x439ce4 PSHUFD $-0x12,%XMM4,%XMM4 |
(312) 0x439ce9 MOVQ %XMM4,%RDX |
(312) 0x439cee MOVDQA %XMM3,%XMM4 |
(312) 0x439cf2 PADDQ %XMM1,%XMM4 |
(312) 0x439cf6 PSLLQ $0x3,%XMM4 |
(312) 0x439cfb PADDQ %XMM11,%XMM4 |
(312) 0x439d00 MOVQ %XMM4,%RAX |
(312) 0x439d05 PSHUFD $-0x12,%XMM4,%XMM4 |
(312) 0x439d0a PADDD 0x51a1e(%RIP),%XMM0 |
(312) 0x439d12 MOVQ %XMM4,%RDI |
(312) 0x439d17 PXOR %XMM4,%XMM4 |
(312) 0x439d1b PCMPGTD %XMM0,%XMM4 |
(312) 0x439d1f PUNPCKLDQ %XMM4,%XMM0 |
(312) 0x439d23 MOVDQA %XMM5,%XMM4 |
(312) 0x439d27 PADDQ %XMM1,%XMM4 |
(312) 0x439d2b PSLLQ $0x3,%XMM4 |
(312) 0x439d30 MOVDQA -0x150(%RBP),%XMM7 |
(312) 0x439d38 PADDQ %XMM7,%XMM4 |
(312) 0x439d3c MOVQ %XMM4,%R14 |
(312) 0x439d41 PSHUFD $-0x12,%XMM4,%XMM4 |
(312) 0x439d46 PADDQ %XMM0,%XMM3 |
(312) 0x439d4a PADDQ %XMM0,%XMM5 |
(312) 0x439d4e MOVQ %XMM4,%R9 |
(312) 0x439d53 PADDQ %XMM0,%XMM14 |
(312) 0x439d58 PADDQ %XMM10,%XMM0 |
(312) 0x439d5d MOVSD (%RSI),%XMM4 |
(312) 0x439d61 PADDQ %XMM1,%XMM10 |
(312) 0x439d66 PSLLQ $0x3,%XMM10 |
(312) 0x439d6c PADDQ %XMM7,%XMM10 |
(312) 0x439d71 MOVQ %XMM10,%RSI |
(312) 0x439d76 PSHUFD $-0x12,%XMM10,%XMM10 |
(312) 0x439d7c PSLLQ $0x3,%XMM3 |
(312) 0x439d81 MOVHPD (%RDX),%XMM4 |
(312) 0x439d85 PADDQ %XMM11,%XMM3 |
(312) 0x439d8a MOVQ %XMM3,%RDX |
(312) 0x439d8f PSHUFD $-0x12,%XMM3,%XMM3 |
(312) 0x439d94 MOVQ %XMM10,%R15 |
(312) 0x439d99 MOVQ %XMM3,%R11 |
(312) 0x439d9e PSLLQ $0x3,%XMM5 |
(312) 0x439da3 PADDQ %XMM7,%XMM5 |
(312) 0x439da7 MOVSD (%RDX),%XMM3 |
(312) 0x439dab MOVQ %XMM5,%RDX |
(312) 0x439db0 PSHUFD $-0x12,%XMM5,%XMM5 |
(312) 0x439db5 MOVQ %XMM5,%RBX |
(312) 0x439dba MOVSD (%RAX),%XMM5 |
(312) 0x439dbe PSLLQ $0x3,%XMM14 |
(312) 0x439dc4 PADDQ %XMM11,%XMM14 |
(312) 0x439dc9 MOVQ %XMM14,%RAX |
(312) 0x439dce MOVSD (%R14),%XMM11 |
(312) 0x439dd3 PSHUFD $-0x12,%XMM14,%XMM10 |
(312) 0x439dd9 MOVQ %XMM10,%R14 |
(312) 0x439dde MOVHPD (%R11),%XMM3 |
(312) 0x439de3 MOVSD (%RSI),%XMM10 |
(312) 0x439de8 PSLLQ $0x3,%XMM0 |
(312) 0x439ded MOVHPD (%RDI),%XMM5 |
(312) 0x439df1 PADDQ %XMM7,%XMM0 |
(312) 0x439df5 MOVQ %XMM0,%RSI |
(312) 0x439dfa MOVHPD (%R9),%XMM11 |
(312) 0x439dff PSHUFD $-0x12,%XMM0,%XMM0 |
(312) 0x439e04 MOVQ %XMM0,%RDI |
(312) 0x439e09 MOVHPD (%R15),%XMM10 |
(312) 0x439e0e MOVSD (%RDX),%XMM14 |
(312) 0x439e13 MOVHPD (%RBX),%XMM14 |
(312) 0x439e18 MULPD %XMM3,%XMM14 |
(312) 0x439e1d MOVSD (%RAX),%XMM3 |
(312) 0x439e21 MOVHPD (%R14),%XMM3 |
(312) 0x439e26 MULPD %XMM5,%XMM11 |
(312) 0x439e2b MOVSD (%RSI),%XMM0 |
(312) 0x439e2f MOVHPD (%RDI),%XMM0 |
(312) 0x439e33 MULPD %XMM4,%XMM10 |
(312) 0x439e38 MULPD %XMM3,%XMM0 |
(312) 0x439e3c MOVDQA -0x130(%RBP),%XMM5 |
(312) 0x439e44 MOVDQA %XMM5,%XMM3 |
(312) 0x439e48 PSRLQ $0x20,%XMM3 |
(312) 0x439e4d PMULUDQ %XMM15,%XMM3 |
(312) 0x439e52 MOVDQA %XMM5,%XMM4 |
(312) 0x439e56 PMULUDQ %XMM8,%XMM4 |
(312) 0x439e5b PADDQ %XMM3,%XMM4 |
(312) 0x439e5f ADDPD %XMM14,%XMM0 |
(312) 0x439e64 MOVDQA %XMM5,%XMM3 |
(312) 0x439e68 PMULUDQ %XMM15,%XMM3 |
(312) 0x439e6d PADDQ %XMM1,%XMM3 |
(312) 0x439e71 PSLLQ $0x20,%XMM4 |
(312) 0x439e76 PADDQ %XMM4,%XMM3 |
(312) 0x439e7a MOVDQA -0xc0(%RBP),%XMM4 |
(312) 0x439e82 MOVDQA %XMM4,%XMM5 |
(312) 0x439e86 PMULUDQ %XMM2,%XMM5 |
(312) 0x439e8a ADDPD %XMM11,%XMM10 |
(312) 0x439e8f MOVDQA %XMM4,%XMM11 |
(312) 0x439e94 PSRLQ $0x20,%XMM11 |
(312) 0x439e9a PMULUDQ %XMM11,%XMM2 |
(312) 0x439e9f PMULUDQ %XMM4,%XMM12 |
(312) 0x439ea4 PADDQ %XMM2,%XMM12 |
(312) 0x439ea9 PSLLQ $0x20,%XMM12 |
(312) 0x439eaf PADDQ %XMM1,%XMM5 |
(312) 0x439eb3 PADDQ %XMM12,%XMM5 |
(312) 0x439eb8 MOVDQA %XMM4,%XMM2 |
(312) 0x439ebc PMULUDQ %XMM15,%XMM2 |
(312) 0x439ec1 PMULUDQ %XMM15,%XMM11 |
(312) 0x439ec6 ADDPD %XMM10,%XMM0 |
(312) 0x439ecb PMULUDQ %XMM8,%XMM4 |
(312) 0x439ed0 PADDQ %XMM11,%XMM4 |
(312) 0x439ed5 PSLLQ $0x20,%XMM4 |
(312) 0x439eda PADDQ %XMM1,%XMM2 |
(312) 0x439ede MOVDQA -0x100(%RBP),%XMM7 |
(312) 0x439ee6 MOVDQA %XMM7,%XMM10 |
(312) 0x439eeb PMULUDQ %XMM15,%XMM10 |
(312) 0x439ef0 PADDQ %XMM4,%XMM2 |
(312) 0x439ef4 MOVDQA %XMM7,%XMM4 |
(312) 0x439ef8 PSRLQ $0x20,%XMM4 |
(312) 0x439efd PMULUDQ %XMM15,%XMM4 |
(312) 0x439f02 PSLLQ $0x3,%XMM3 |
(312) 0x439f07 PADDQ -0x120(%RBP),%XMM3 |
(312) 0x439f0f MOVQ %XMM3,%RAX |
(312) 0x439f14 PSHUFD $-0x12,%XMM3,%XMM3 |
(312) 0x439f19 MOVQ %XMM3,%RDX |
(312) 0x439f1e PMULUDQ %XMM7,%XMM8 |
(312) 0x439f23 PADDQ %XMM4,%XMM8 |
(312) 0x439f28 PADDQ %XMM1,%XMM10 |
(312) 0x439f2d PSLLQ $0x3,%XMM5 |
(312) 0x439f32 MOVDQA -0x110(%RBP),%XMM3 |
(312) 0x439f3a PADDQ %XMM3,%XMM5 |
(312) 0x439f3e MOVQ %XMM5,%RSI |
(312) 0x439f43 PSHUFD $-0x12,%XMM5,%XMM1 |
(312) 0x439f48 MOVQ %XMM1,%RDI |
(312) 0x439f4d MULPD 0x5164b(%RIP),%XMM0 |
(312) 0x439f55 MOVLPD %XMM0,(%RAX) |
(312) 0x439f59 PSLLQ $0x3,%XMM2 |
(312) 0x439f5e MOVHPD %XMM0,(%RDX) |
(312) 0x439f62 PADDQ %XMM3,%XMM2 |
(312) 0x439f66 MOVQ %XMM2,%RAX |
(312) 0x439f6b PSHUFD $-0x12,%XMM2,%XMM1 |
(312) 0x439f70 MOVQ %XMM1,%RDX |
(312) 0x439f75 MOVSD (%RSI),%XMM1 |
(312) 0x439f79 MOVHPD (%RDI),%XMM1 |
(312) 0x439f7d SUBPD %XMM1,%XMM0 |
(312) 0x439f81 MOVSD (%RAX),%XMM1 |
(312) 0x439f85 MOVHPD (%RDX),%XMM1 |
(312) 0x439f89 ADDPD %XMM0,%XMM1 |
(312) 0x439f8d PSLLQ $0x20,%XMM8 |
(312) 0x439f93 PADDQ %XMM8,%XMM10 |
(312) 0x439f98 PSLLQ $0x3,%XMM10 |
(312) 0x439f9e PADDQ -0xf0(%RBP),%XMM10 |
(312) 0x439fa7 MOVQ %XMM10,%RAX |
(312) 0x439fac PSHUFD $-0x12,%XMM10,%XMM0 |
(312) 0x439fb2 MOVQ %XMM0,%RDX |
(312) 0x439fb7 MOVLPD %XMM1,(%RAX) |
(312) 0x439fbb MOVHPD %XMM1,(%RDX) |
(312) 0x439fbf PADDQ 0x515e8(%RIP),%XMM13 |
(312) 0x439fc8 ADD $0x2,%R8 |
(312) 0x439fcc CMP %R12,%R8 |
(312) 0x439fcf JB 439b30 |
0x439fd5 CMP %R12,-0x68(%RBP) |
0x439fd9 MOV -0x80(%RBP),%RBX |
0x439fdd MOV -0x48(%RBP),%R11 |
0x439fe1 MOV -0x78(%RBP),%R14 |
0x439fe5 MOV -0x70(%RBP),%R15 |
0x439fe9 JNE 43a026 |
0x439feb MOV $0x4aa920,%EDI |
0x439ff0 MOV -0x3c(%RBP),%ESI |
0x439ff3 ADD $0x168,%RSP |
0x439ffa POP %RBX |
0x439ffb POP %R12 |
0x439ffd POP %R13 |
0x439fff POP %R14 |
0x43a001 POP %R15 |
0x43a003 POP %RBP |
0x43a004 JMP 404050 |
0x43a009 ADD $0x168,%RSP |
0x43a010 POP %RBX |
0x43a011 POP %R12 |
0x43a013 POP %R13 |
0x43a015 POP %R14 |
0x43a017 POP %R15 |
0x43a019 POP %RBP |
0x43a01a RET |
0x43a01b MOV %RDX,%R14 |
0x43a01e MOV %RSI,%R15 |
0x43a021 JMP 43a142 |
0x43a026 ADD %R12,%RCX |
0x43a029 JMP 43a142 |
0x43a02e XCHG %AX,%AX |
(311) 0x43a030 MOV %RCX,%RAX |
(311) 0x43a033 CQTO |
(311) 0x43a035 IDIV %RBX |
(311) 0x43a038 ADD %R11D,%EDX |
(311) 0x43a03b MOVSXD %EDX,%RAX |
(311) 0x43a03e LEA -0x1(%RSI),%EDX |
(311) 0x43a041 MOVSXD %EDX,%RDI |
(311) 0x43a044 MOV %R15,%R8 |
(311) 0x43a047 IMUL %RDI,%R8 |
(311) 0x43a04b MOVSXD %ESI,%RDX |
(311) 0x43a04e MOV %R15,%RSI |
(311) 0x43a051 IMUL %RDX,%RSI |
(311) 0x43a055 LEA (%R8,%RAX,1),%R9 |
(311) 0x43a059 DEC %R9 |
(311) 0x43a05c ADD %RAX,%R8 |
(311) 0x43a05f MOVSD (%R13,%R8,8),%XMM0 |
(311) 0x43a066 LEA (%RSI,%RAX,1),%R8 |
(311) 0x43a06a DEC %R8 |
(311) 0x43a06d ADD %RAX,%RSI |
(311) 0x43a070 MOVHPD (%R13,%RSI,8),%XMM0 |
(311) 0x43a077 MOV %R14,%RSI |
(311) 0x43a07a IMUL %RDI,%RSI |
(311) 0x43a07e MOVSD (%R13,%R9,8),%XMM1 |
(311) 0x43a085 MOV %R14,%R9 |
(311) 0x43a088 IMUL %RDX,%R9 |
(311) 0x43a08c MOVHPD (%R13,%R8,8),%XMM1 |
(311) 0x43a093 LEA (%RSI,%RAX,1),%R8 |
(311) 0x43a097 DEC %R8 |
(311) 0x43a09a ADD %RAX,%RSI |
(311) 0x43a09d MOVSD (%R10,%RSI,8),%XMM2 |
(311) 0x43a0a3 LEA (%R9,%RAX,1),%RSI |
(311) 0x43a0a7 DEC %RSI |
(311) 0x43a0aa ADD %RAX,%R9 |
(311) 0x43a0ad MOVHPD (%R10,%R9,8),%XMM2 |
(311) 0x43a0b3 MOVSD (%R10,%R8,8),%XMM3 |
(311) 0x43a0b9 MULPD %XMM0,%XMM2 |
(311) 0x43a0bd MOVHPD (%R10,%RSI,8),%XMM3 |
(311) 0x43a0c3 MULPD %XMM1,%XMM3 |
(311) 0x43a0c7 ADDPD %XMM2,%XMM3 |
(311) 0x43a0cb MOVAPD %XMM3,%XMM0 |
(311) 0x43a0cf UNPCKHPD %XMM3,%XMM0 |
(311) 0x43a0d3 ADDSD %XMM3,%XMM0 |
(311) 0x43a0d7 MULSD 0x514a9(%RIP),%XMM0 |
(311) 0x43a0df MOV -0xa0(%RBP),%RSI |
(311) 0x43a0e6 IMUL %RDX,%RSI |
(311) 0x43a0ea ADD %RAX,%RSI |
(311) 0x43a0ed MOV -0x98(%RBP),%R8 |
(311) 0x43a0f4 MOVSD %XMM0,(%R8,%RSI,8) |
(311) 0x43a0fa IMUL %R12,%RDI |
(311) 0x43a0fe ADD %RAX,%RDI |
(311) 0x43a101 MOV -0x90(%RBP),%R8 |
(311) 0x43a108 SUBSD (%R8,%RDI,8),%XMM0 |
(311) 0x43a10e MOV %R12,%RSI |
(311) 0x43a111 IMUL %RDX,%RSI |
(311) 0x43a115 ADD %RAX,%RSI |
(311) 0x43a118 ADDSD (%R8,%RSI,8),%XMM0 |
(311) 0x43a11e IMUL -0x88(%RBP),%RDX |
(311) 0x43a126 ADD %RAX,%RDX |
(311) 0x43a129 MOV -0x38(%RBP),%RAX |
(311) 0x43a12d MOVSD %XMM0,(%RAX,%RDX,8) |
(311) 0x43a132 INC %RCX |
(311) 0x43a135 CMP %RCX,-0xa8(%RBP) |
(311) 0x43a13c JE 439feb |
(311) 0x43a142 MOV %RCX,%RDI |
(311) 0x43a145 SHR $0x20,%RDI |
(311) 0x43a149 JE 43a170 |
(311) 0x43a14b MOV %RCX,%RAX |
(311) 0x43a14e XOR %EDX,%EDX |
(311) 0x43a150 DIV %RBX |
(311) 0x43a153 MOV %RAX,%RSI |
(311) 0x43a156 MOV -0x50(%RBP),%R12 |
(311) 0x43a15a ADD -0x2c(%RBP),%ESI |
(311) 0x43a15d TEST %RDI,%RDI |
(311) 0x43a160 JNE 43a030 |
(311) 0x43a166 JMP 43a188 |
0x43a168 NOPL (%RAX,%RAX,1) |
(311) 0x43a170 MOV %ECX,%EAX |
(311) 0x43a172 XOR %EDX,%EDX |
(311) 0x43a174 DIV %EBX |
(311) 0x43a176 MOV %EAX,%ESI |
(311) 0x43a178 MOV -0x50(%RBP),%R12 |
(311) 0x43a17c ADD -0x2c(%RBP),%ESI |
(311) 0x43a17f TEST %RDI,%RDI |
(311) 0x43a182 JNE 43a030 |
(311) 0x43a188 MOV %ECX,%EAX |
(311) 0x43a18a XOR %EDX,%EDX |
(311) 0x43a18c DIV %EBX |
(311) 0x43a18e JMP 43a038 |
0x43a193 NOPW %CS:(%RAX,%RAX,1) |
Path / |
Source file and lines | advec_mom.cpp:167-172 |
Module | exec |
nb instructions | 152 |
nb uops | 159 |
loop length | 688 |
used x86 registers | 16 |
used mmx registers | 0 |
used xmm registers | 4 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 37 |
micro-operation queue | 26.50 cycles |
front end | 26.50 cycles |
ALU0/BRU0 | ALU1 | ALU2 | ALU3 | BRU1 | AGU0 | AGU1 | AGU2 | FP0 | FP1 | FP2 | FP3 | FP4 | FP5 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 7.00 | 7.00 | 6.75 | 6.75 | 4.50 | 21.00 | 21.00 | 21.00 | 6.50 | 6.50 | 6.50 | 6.50 | 7.00 | 7.00 |
cycles | 7.00 | 7.00 | 6.75 | 6.75 | 4.50 | 21.00 | 21.00 | 21.00 | 6.50 | 6.50 | 6.50 | 6.50 | 7.00 | 7.00 |
Cycles executing div or sqrt instructions | NA |
Front-end | 26.50 |
Dispatch | 21.00 |
Overall L1 | 26.50 |
all | 35% |
load | 5% |
store | 42% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 14% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 48% |
all | 15% |
load | 12% |
store | 16% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 13% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 17% |
Instruction | Nb FU | ALU0/BRU0 | ALU1 | ALU2 | ALU3 | BRU1 | AGU0 | AGU1 | AGU2 | FP0 | FP1 | FP2 | FP3 | FP4 | FP5 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
SUB $0x168,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV %RCX,%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x38(%RBP),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x28(%RBP),%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x20(%RBP),%R10 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%RBP),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x18(%RBP),%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV %ECX,-0x2c(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVL $0,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
TEST %RAX,%RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
JS 43a009 <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x719> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 |
MOV %R8,%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RDX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R9,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOVQ $0,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %RAX,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVQ $0x1,-0xb0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
SUB $0x8,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
LEA -0xb0(%RBP),%RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
LEA -0x40(%RBP),%RCX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
LEA -0x60(%RBP),%R8 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
LEA -0x58(%RBP),%R9 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV $0x4aa900,%EDI | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV %ESI,-0x3c(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV $0x22,%EDX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
PUSH $0x1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH $0x1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %R10,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
CALL 404240 <__kmpc_for_static_init_8@plt> | 2 | 0.50 | 0 | 0 | 0 | 0.50 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x48(%RBP),%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
ADD $0x20,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV -0x60(%RBP),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x58(%RBP),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
CMP %RDX,%RCX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
JA 439feb <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x6fb> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 |
MOV %RDX,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
SUB %R11D,%EBX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV (%R13),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R13),%R10 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV (%R14),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R14),%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV (%R12),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R12),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV (%R15),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R15),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x38(%RBP),%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV (%R12),%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R12),%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV %R12,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
INC %RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV %RAX,-0xa8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
SUB %RCX,%RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV $-0x2,%R12D | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
AND %RAX,%R12 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV %R9,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %R14,-0x50(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %RDI,-0xa0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %R8,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %R15,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
JE 43a01b <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x72b> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 |
MOV %RAX,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %RBX,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVD %EBX,%XMM0 | 1 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2 | 1 |
PSHUFD $0x44,%XMM0,%XMM1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVD -0x2c(%RBP),%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.50 |
PSHUFD $0x50,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x180(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVD %R11D,%XMM0 | 1 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2 | 1 |
PSHUFD $0x50,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x170(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOV %RDX,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVQ %RDX,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x140(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R10,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x160(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOV %RSI,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVQ %RSI,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0xd0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R13,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x150(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %RDI,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x130(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R8,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x120(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R14,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R9,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x110(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R15,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x100(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ -0x38(%RBP),%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0xf0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %RCX,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
PADDQ 0x51a7e(%RIP),%XMM13 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 1 | 0.50 |
XOR %R8D,%R8D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 |
MOVDQA %XMM1,-0x190(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
PSHUFD $-0x12,%XMM1,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0xe0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
PXOR %XMM6,%XMM6 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 |
XCHG %AX,%AX | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 |
CMP %R12,-0x68(%RBP) | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
MOV -0x80(%RBP),%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x48(%RBP),%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x78(%RBP),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x70(%RBP),%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
JNE 43a026 <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x736> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 |
MOV $0x4aa920,%EDI | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV -0x3c(%RBP),%ESI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
ADD $0x168,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
POP %RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
JMP 404050 <__kmpc_for_static_fini@plt> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
ADD $0x168,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
POP %RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
RET | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %RDX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RSI,%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 43a142 <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x852> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
ADD %R12,%RCX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
JMP 43a142 <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x852> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
XCHG %AX,%AX | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 |
NOPL (%RAX,%RAX,1) | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 |
NOPW %CS:(%RAX,%RAX,1) | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 |
Source file and lines | advec_mom.cpp:167-172 |
Module | exec |
nb instructions | 152 |
nb uops | 159 |
loop length | 688 |
used x86 registers | 16 |
used mmx registers | 0 |
used xmm registers | 4 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 37 |
micro-operation queue | 26.50 cycles |
front end | 26.50 cycles |
ALU0/BRU0 | ALU1 | ALU2 | ALU3 | BRU1 | AGU0 | AGU1 | AGU2 | FP0 | FP1 | FP2 | FP3 | FP4 | FP5 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 7.00 | 7.00 | 6.75 | 6.75 | 4.50 | 21.00 | 21.00 | 21.00 | 6.50 | 6.50 | 6.50 | 6.50 | 7.00 | 7.00 |
cycles | 7.00 | 7.00 | 6.75 | 6.75 | 4.50 | 21.00 | 21.00 | 21.00 | 6.50 | 6.50 | 6.50 | 6.50 | 7.00 | 7.00 |
Cycles executing div or sqrt instructions | NA |
Front-end | 26.50 |
Dispatch | 21.00 |
Overall L1 | 26.50 |
all | 35% |
load | 5% |
store | 42% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 14% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 48% |
all | 15% |
load | 12% |
store | 16% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 13% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 17% |
Instruction | Nb FU | ALU0/BRU0 | ALU1 | ALU2 | ALU3 | BRU1 | AGU0 | AGU1 | AGU2 | FP0 | FP1 | FP2 | FP3 | FP4 | FP5 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
SUB $0x168,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV %RCX,%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x38(%RBP),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x28(%RBP),%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x20(%RBP),%R10 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%RBP),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x18(%RBP),%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV %ECX,-0x2c(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVL $0,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
TEST %RAX,%RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
JS 43a009 <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x719> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 |
MOV %R8,%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RDX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R9,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOVQ $0,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %RAX,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVQ $0x1,-0xb0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
SUB $0x8,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
LEA -0xb0(%RBP),%RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
LEA -0x40(%RBP),%RCX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
LEA -0x60(%RBP),%R8 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
LEA -0x58(%RBP),%R9 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV $0x4aa900,%EDI | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV %ESI,-0x3c(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV $0x22,%EDX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
PUSH $0x1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH $0x1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
PUSH %RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %R10,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
CALL 404240 <__kmpc_for_static_init_8@plt> | 2 | 0.50 | 0 | 0 | 0 | 0.50 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x48(%RBP),%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
ADD $0x20,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV -0x60(%RBP),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x58(%RBP),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
CMP %RDX,%RCX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
JA 439feb <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x6fb> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 |
MOV %RDX,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
SUB %R11D,%EBX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV (%R13),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R13),%R10 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV (%R14),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R14),%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV (%R12),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R12),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV (%R15),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R15),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x38(%RBP),%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV (%R12),%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV 0x10(%R12),%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV %R12,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
INC %RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV %RAX,-0xa8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
SUB %RCX,%RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV $-0x2,%R12D | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
AND %RAX,%R12 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV %R9,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %R14,-0x50(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %RDI,-0xa0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %R8,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %R15,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
JE 43a01b <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x72b> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 |
MOV %RAX,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %RBX,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVD %EBX,%XMM0 | 1 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2 | 1 |
PSHUFD $0x44,%XMM0,%XMM1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVD -0x2c(%RBP),%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.50 |
PSHUFD $0x50,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x180(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVD %R11D,%XMM0 | 1 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2 | 1 |
PSHUFD $0x50,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x170(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOV %RDX,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVQ %RDX,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x140(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R10,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x160(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOV %RSI,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOVQ %RSI,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0xd0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R13,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x150(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %RDI,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x130(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R8,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x120(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R14,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R9,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x110(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %R15,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0x100(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ -0x38(%RBP),%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
PSHUFD $0x44,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0xf0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
MOVQ %RCX,%XMM0 | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 6 | 1 |
PSHUFD $0x44,%XMM0,%XMM13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
PADDQ 0x51a7e(%RIP),%XMM13 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 1 | 0.50 |
XOR %R8D,%R8D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 |
MOVDQA %XMM1,-0x190(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
PSHUFD $-0x12,%XMM1,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 1 | 0.33 |
MOVDQA %XMM0,-0xe0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0.50 | 0.50 | 4 | 1 |
PXOR %XMM6,%XMM6 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.25 |
XCHG %AX,%AX | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 |
CMP %R12,-0x68(%RBP) | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
MOV -0x80(%RBP),%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x48(%RBP),%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x78(%RBP),%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
MOV -0x70(%RBP),%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
JNE 43a026 <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x736> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 |
MOV $0x4aa920,%EDI | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
MOV -0x3c(%RBP),%ESI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 |
ADD $0x168,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
POP %RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
JMP 404050 <__kmpc_for_static_fini@plt> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
ADD $0x168,%RSP | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
POP %RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 |
RET | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %RDX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RSI,%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 43a142 <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x852> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
ADD %R12,%RCX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 |
JMP 43a142 <_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12+0x852> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
XCHG %AX,%AX | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 |
NOPL (%RAX,%RAX,1) | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 |
NOPW %CS:(%RAX,%RAX,1) | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 |
Name | Coverage (%) | Time (s) |
---|---|---|
▼_Z16advec_mom_kerneliiiiRN6clover8Buffer2DIdEES2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_S2_RNS_8Buffer1DIdEES5_iii.extracted.12– | 2.81 | 1.54 |
○Loop 312 - advec_mom.cpp:167-172 - exec | 2.81 | 1.54 |
○Loop 311 - advec_mom.cpp:167-172 - exec | 0 | 0 |