Function: field_summary_kernel_.DIR.OMP.PARALLEL.2 | Module: exec | Source: field_summary_kernel.f90:54-74 | Coverage: 0.29% |
---|
Function: field_summary_kernel_.DIR.OMP.PARALLEL.2 | Module: exec | Source: field_summary_kernel.f90:54-74 | Coverage: 0.29% |
---|
/scratch_na/users/xoserete/qaas_runs/171-215-0463/intel/CloverLeafFC/build/CloverLeafFC/CloverLeaf_ref/kernels/field_summary_kernel.f90: 54 - 74 |
-------------------------------------------------------------------------------- |
54: !$OMP PARALLEL |
55: !$OMP DO PRIVATE(vsqrd,cell_vol,cell_mass) REDUCTION(+ : vol,mass,press,ie,ke) |
56: DO k=y_min,y_max |
57: !$OMP SIMD |
58: DO j=x_min,x_max |
59: vsqrd=0.0 |
60: DO kv=k,k+1 |
61: DO jv=j,j+1 |
62: vsqrd=vsqrd+0.25*(xvel0(jv,kv)**2+yvel0(jv,kv)**2) |
63: ENDDO |
64: ENDDO |
65: cell_vol=volume(j,k) |
66: cell_mass=cell_vol*density0(j,k) |
67: vol=vol+cell_vol |
68: mass=mass+cell_mass |
69: ie=ie+cell_mass*energy0(j,k) |
70: ke=ke+cell_mass*0.5*vsqrd |
71: press=press+cell_vol*pressure(j,k) |
72: ENDDO |
73: ENDDO |
74: !$OMP END DO |
0x43c430 PUSH %RBP |
0x43c431 MOV %RSP,%RBP |
0x43c434 PUSH %R15 |
0x43c436 PUSH %R14 |
0x43c438 PUSH %R13 |
0x43c43a PUSH %R12 |
0x43c43c PUSH %RBX |
0x43c43d SUB $0xb8,%RSP |
0x43c444 MOV %R9,-0x90(%RBP) |
0x43c44b MOV %R8,-0x88(%RBP) |
0x43c452 MOV 0x60(%RBP),%EBX |
0x43c455 MOV 0x58(%RBP),%EAX |
0x43c458 SUB %EBX,%EAX |
0x43c45a MOVL $0,-0x58(%RBP) |
0x43c461 JS 43cd61 |
0x43c467 MOV %RCX,%R14 |
0x43c46a MOV %RDX,%R13 |
0x43c46d MOV %RDI,-0x30(%RBP) |
0x43c471 MOV (%RDI),%ESI |
0x43c473 MOVL $0,-0x38(%RBP) |
0x43c47a MOV %EAX,-0x34(%RBP) |
0x43c47d MOVL $0x1,-0x54(%RBP) |
0x43c484 SUB $0x8,%RSP |
0x43c488 LEA -0x54(%RBP),%RAX |
0x43c48c LEA -0x58(%RBP),%RCX |
0x43c490 LEA -0x38(%RBP),%R8 |
0x43c494 LEA -0x34(%RBP),%R9 |
0x43c498 MOV $0x74a5b0,%EDI |
0x43c49d MOV %ESI,-0x4c(%RBP) |
0x43c4a0 MOV $0x22,%EDX |
0x43c4a5 PUSH $0x1 |
0x43c4a7 PUSH $0x1 |
0x43c4a9 PUSH %RAX |
0x43c4aa CALL 4045a0 <__kmpc_for_static_init_4@plt> |
0x43c4af ADD $0x20,%RSP |
0x43c4b3 MOV -0x38(%RBP),%EAX |
0x43c4b6 MOV -0x34(%RBP),%R9D |
0x43c4ba VXORPD %XMM0,%XMM0,%XMM0 |
0x43c4be VXORPD %XMM2,%XMM2,%XMM2 |
0x43c4c2 VXORPD %XMM3,%XMM3,%XMM3 |
0x43c4c6 VXORPD %XMM4,%XMM4,%XMM4 |
0x43c4ca VXORPD %XMM1,%XMM1,%XMM1 |
0x43c4ce MOV %RAX,-0x40(%RBP) |
0x43c4d2 SUB %EAX,%R9D |
0x43c4d5 JAE 43c600 |
0x43c4db MOV 0x20(%RBP),%R15 |
0x43c4df MOV 0x18(%RBP),%R14 |
0x43c4e3 MOV 0x10(%RBP),%R13 |
0x43c4e7 VMOVSD %XMM1,-0x80(%RBP) |
0x43c4ec VMOVSD %XMM4,-0x78(%RBP) |
0x43c4f1 VMOVSD %XMM3,-0x70(%RBP) |
0x43c4f6 VMOVSD %XMM2,-0x68(%RBP) |
0x43c4fb VMOVSD %XMM0,-0x60(%RBP) |
0x43c500 MOV $0x74a5d0,%EDI |
0x43c505 MOV -0x4c(%RBP),%ESI |
0x43c508 VZEROUPPER |
0x43c50b CALL 404190 <__kmpc_for_static_fini@plt> |
0x43c510 MOV -0x30(%RBP),%RAX |
0x43c514 MOV (%RAX),%ESI |
0x43c516 SUB $0x8,%RSP |
0x43c51a MOV $0x7a3060,%RAX |
0x43c521 LEA -0x80(%RBP),%R8 |
0x43c525 MOV $0x74a690,%EDI |
0x43c52a MOV $0x43c3e0,%R9D |
0x43c530 MOV $0x5,%EDX |
0x43c535 MOV $0x28,%ECX |
0x43c53a PUSH %RAX |
0x43c53b CALL 404780 <__kmpc_reduce@plt> |
0x43c540 ADD $0x10,%RSP |
0x43c544 CMP $0x2,%EAX |
0x43c547 JGE 43ccc0 |
0x43c54d CMP $0x1,%EAX |
0x43c550 MOV -0x30(%RBP),%RDI |
0x43c554 JNE 43cd61 |
0x43c55a VMOVSD -0x80(%RBP),%XMM0 |
0x43c55f VADDSD (%R15),%XMM0,%XMM0 |
0x43c564 VMOVSD %XMM0,(%R15) |
0x43c569 VMOVSD -0x78(%RBP),%XMM0 |
0x43c56e VADDSD (%R14),%XMM0,%XMM0 |
0x43c573 VMOVSD %XMM0,(%R14) |
0x43c578 VMOVSD -0x70(%RBP),%XMM0 |
0x43c57d MOV -0x88(%RBP),%RAX |
0x43c584 VADDSD (%RAX),%XMM0,%XMM0 |
0x43c588 VMOVSD %XMM0,(%RAX) |
0x43c58c VMOVSD -0x68(%RBP),%XMM0 |
0x43c591 VADDSD (%R13),%XMM0,%XMM0 |
0x43c597 VMOVSD %XMM0,(%R13) |
0x43c59d VMOVSD -0x60(%RBP),%XMM0 |
0x43c5a2 MOV -0x90(%RBP),%RAX |
0x43c5a9 VADDSD (%RAX),%XMM0,%XMM0 |
0x43c5ad VMOVSD %XMM0,(%RAX) |
0x43c5b1 MOV (%RDI),%ESI |
0x43c5b3 MOV $0x7a3060,%RDX |
0x43c5ba MOV $0x74a6b0,%EDI |
0x43c5bf JMP 43cd58 |
0x43c5c4 NOPW %CS:(%RAX,%RAX,1) |
0x43c5d3 NOPW %CS:(%RAX,%RAX,1) |
0x43c5e2 NOPW %CS:(%RAX,%RAX,1) |
0x43c5f1 NOPW %CS:(%RAX,%RAX,1) |
0x43c600 MOV 0xa0(%RBP),%RDI |
0x43c607 MOV 0x98(%RBP),%RSI |
0x43c60e MOV 0x90(%RBP),%R15 |
0x43c615 MOV 0x88(%RBP),%R12 |
0x43c61c MOV 0x70(%RBP),%RAX |
0x43c620 MOV 0x68(%RBP),%RCX |
0x43c624 ADD $-0x2,%R14D |
0x43c628 MOVSXD %R14D,%R14 |
0x43c62b ADD $-0x2,%R13D |
0x43c62f MOVSXD %R13D,%RDX |
0x43c632 MOV %RDX,-0x98(%RBP) |
0x43c639 MOV -0x40(%RBP),%RDX |
0x43c63d ADD %EBX,%EDX |
0x43c63f MOV %RDX,-0x40(%RBP) |
0x43c643 MOVSXD (%RAX),%RDX |
0x43c646 MOV (%RCX),%EAX |
0x43c648 SUB %EDX,%EAX |
0x43c64a MOV %RAX,-0xc0(%RBP) |
0x43c651 INC %EAX |
0x43c653 CMP $0x2,%EAX |
0x43c656 MOV $0x1,%ECX |
0x43c65b CMOVGE %EAX,%ECX |
0x43c65e MOV %RDX,-0xe0(%RBP) |
0x43c665 VPBROADCASTD %EDX,%YMM0 |
0x43c66b MOV %RCX,-0x48(%RBP) |
0x43c66f MOV %ECX,%R8D |
0x43c672 AND $0x7ffffff8,%R8D |
0x43c679 VPMOVSXDQ %YMM0,%ZMM0 |
0x43c67f VMOVDQA64 0xcf777(%RIP),%ZMM5 |
0x43c689 VPBROADCASTQ 0xcfae5(%RIP),%ZMM6 |
0x43c693 VBROADCASTSD 0xcfad3(%RIP),%ZMM7 |
0x43c69d VMOVDQA %XMM0,%XMM8 |
0x43c6a1 VPXOR %XMM0,%XMM0,%XMM0 |
0x43c6a5 XOR %R11D,%R11D |
0x43c6a8 MOV %R9D,-0x50(%RBP) |
0x43c6ac MOV %R8,-0xa8(%RBP) |
0x43c6b3 JMP 43c7cf |
0x43c6b8 NOPL (%RAX,%RAX,1) |
(237) 0x43c6c0 VEXTRACTF64X4 $0x1,%ZMM14,%YMM10 |
(237) 0x43c6c7 VADDPD %ZMM10,%ZMM14,%ZMM10 |
(237) 0x43c6cd VMOVAPD %XMM10,%XMM14 |
(237) 0x43c6d2 VEXTRACTF128 $0x1,%YMM10,%XMM10 |
(237) 0x43c6d8 VADDPD %XMM10,%XMM14,%XMM10 |
(237) 0x43c6dd VSHUFPD $0x1,%XMM10,%XMM10,%XMM14 |
(237) 0x43c6e3 VADDSD %XMM14,%XMM10,%XMM10 |
(237) 0x43c6e8 VADDSD %XMM3,%XMM10,%XMM3 |
(237) 0x43c6ec VEXTRACTF64X4 $0x1,%ZMM13,%YMM10 |
(237) 0x43c6f3 VADDPD %ZMM10,%ZMM13,%ZMM10 |
(237) 0x43c6f9 VMOVAPD %XMM10,%XMM13 |
(237) 0x43c6fe VEXTRACTF128 $0x1,%YMM10,%XMM10 |
(237) 0x43c704 VADDPD %XMM10,%XMM13,%XMM10 |
(237) 0x43c709 VSHUFPD $0x1,%XMM10,%XMM10,%XMM13 |
(237) 0x43c70f VADDSD %XMM13,%XMM10,%XMM10 |
(237) 0x43c714 VADDSD %XMM0,%XMM10,%XMM0 |
(237) 0x43c718 VEXTRACTF64X4 $0x1,%ZMM12,%YMM10 |
(237) 0x43c71f VADDPD %ZMM10,%ZMM12,%ZMM10 |
(237) 0x43c725 VMOVAPD %XMM10,%XMM12 |
(237) 0x43c72a VEXTRACTF128 $0x1,%YMM10,%XMM10 |
(237) 0x43c730 VADDPD %XMM10,%XMM12,%XMM10 |
(237) 0x43c735 VSHUFPD $0x1,%XMM10,%XMM10,%XMM12 |
(237) 0x43c73b VADDSD %XMM12,%XMM10,%XMM10 |
(237) 0x43c740 VADDSD %XMM2,%XMM10,%XMM2 |
(237) 0x43c744 VEXTRACTF64X4 $0x1,%ZMM11,%YMM10 |
(237) 0x43c74b VADDPD %ZMM10,%ZMM11,%ZMM10 |
(237) 0x43c751 VMOVAPD %XMM10,%XMM11 |
(237) 0x43c756 VEXTRACTF128 $0x1,%YMM10,%XMM10 |
(237) 0x43c75c VADDPD %XMM10,%XMM11,%XMM10 |
(237) 0x43c761 VSHUFPD $0x1,%XMM10,%XMM10,%XMM11 |
(237) 0x43c767 VADDSD %XMM11,%XMM10,%XMM10 |
(237) 0x43c76c VADDSD %XMM4,%XMM10,%XMM4 |
(237) 0x43c770 VEXTRACTF64X4 $0x1,%ZMM9,%YMM10 |
(237) 0x43c777 VADDPD %ZMM10,%ZMM9,%ZMM9 |
(237) 0x43c77d VMOVAPD %XMM9,%XMM10 |
(237) 0x43c782 VEXTRACTF128 $0x1,%YMM9,%XMM9 |
(237) 0x43c788 VADDPD %XMM9,%XMM10,%XMM9 |
(237) 0x43c78d VSHUFPD $0x1,%XMM9,%XMM9,%XMM10 |
(237) 0x43c793 VADDSD %XMM10,%XMM9,%XMM9 |
(237) 0x43c798 VADDSD %XMM1,%XMM9,%XMM1 |
(237) 0x43c79c MOV -0xa8(%RBP),%R8 |
(237) 0x43c7a3 MOV 0xa0(%RBP),%RDI |
(237) 0x43c7aa MOV 0x98(%RBP),%RSI |
(237) 0x43c7b1 MOV 0x90(%RBP),%R15 |
(237) 0x43c7b8 MOV 0x88(%RBP),%R12 |
(237) 0x43c7bf LEA 0x1(%R11),%EAX |
(237) 0x43c7c3 CMP %R9D,%R11D |
(237) 0x43c7c6 MOV %EAX,%R11D |
(237) 0x43c7c9 JE 43c4db |
(237) 0x43c7cf CMPL $0,-0xc0(%RBP) |
(237) 0x43c7d6 JS 43c7bf |
(237) 0x43c7d8 MOV -0x40(%RBP),%RAX |
(237) 0x43c7dc ADD %R11D,%EAX |
(237) 0x43c7df MOV 0x78(%RBP),%RCX |
(237) 0x43c7e3 MOV (%RCX),%R13 |
(237) 0x43c7e6 MOV 0x80(%RBP),%RCX |
(237) 0x43c7ed MOV (%RCX),%RCX |
(237) 0x43c7f0 MOV (%R12),%RBX |
(237) 0x43c7f4 MOV (%R15),%R15 |
(237) 0x43c7f7 MOV (%RSI),%RSI |
(237) 0x43c7fa MOV (%RDI),%R10 |
(237) 0x43c7fd MOVSXD %EAX,%RDX |
(237) 0x43c800 CMPL $0x8,-0x48(%RBP) |
(237) 0x43c804 MOV %R15,-0xd8(%RBP) |
(237) 0x43c80b MOV %RSI,-0xa0(%RBP) |
(237) 0x43c812 MOV %R10,-0xd0(%RBP) |
(237) 0x43c819 MOV %RDX,-0xc8(%RBP) |
(237) 0x43c820 JAE 43c840 |
(237) 0x43c822 XOR %EDX,%EDX |
(237) 0x43c824 JMP 43cacb |
0x43c829 NOPW %CS:(%RAX,%RAX,1) |
0x43c838 NOPL (%RAX,%RAX,1) |
(237) 0x43c840 MOV %R11,-0xb8(%RBP) |
(237) 0x43c847 SUB -0x98(%RBP),%RDX |
(237) 0x43c84e MOV %RBX,%R9 |
(237) 0x43c851 MOV %R8,%RBX |
(237) 0x43c854 MOV %R13,%RDI |
(237) 0x43c857 IMUL %RDX,%RDI |
(237) 0x43c85b MOV 0x30(%RBP),%RAX |
(237) 0x43c85f ADD %RAX,%RDI |
(237) 0x43c862 MOV %RCX,%R8 |
(237) 0x43c865 IMUL %RDX,%R8 |
(237) 0x43c869 MOV 0x28(%RBP),%RSI |
(237) 0x43c86d ADD %RSI,%R8 |
(237) 0x43c870 LEA 0x1(%RDX),%R11 |
(237) 0x43c874 MOV %R13,%R12 |
(237) 0x43c877 IMUL %R11,%R12 |
(237) 0x43c87b ADD %RAX,%R12 |
(237) 0x43c87e IMUL %RCX,%R11 |
(237) 0x43c882 ADD %RSI,%R11 |
(237) 0x43c885 MOV %R9,-0xb0(%RBP) |
(237) 0x43c88c IMUL %RDX,%R9 |
(237) 0x43c890 ADD 0x50(%RBP),%R9 |
(237) 0x43c894 MOV %R15,%RAX |
(237) 0x43c897 IMUL %RDX,%RAX |
(237) 0x43c89b ADD 0x48(%RBP),%RAX |
(237) 0x43c89f MOV -0xa0(%RBP),%R15 |
(237) 0x43c8a6 IMUL %RDX,%R15 |
(237) 0x43c8aa ADD 0x40(%RBP),%R15 |
(237) 0x43c8ae IMUL %R10,%RDX |
(237) 0x43c8b2 ADD 0x38(%RBP),%RDX |
(237) 0x43c8b6 VXORPD %XMM13,%XMM13,%XMM13 |
(237) 0x43c8bb VXORPD %XMM12,%XMM12,%XMM12 |
(237) 0x43c8c0 VXORPD %XMM11,%XMM11,%XMM11 |
(237) 0x43c8c5 VXORPD %XMM10,%XMM10,%XMM10 |
(237) 0x43c8ca VXORPD %XMM9,%XMM9,%XMM9 |
(237) 0x43c8cf XOR %R10D,%R10D |
(237) 0x43c8d2 VMOVDQA64 %ZMM5,%ZMM14 |
(237) 0x43c8d8 NOPL (%RAX,%RAX,1) |
(238) 0x43c8e0 VMOVDQA %XMM14,%XMM15 |
(238) 0x43c8e5 VPADDQ %XMM8,%XMM14,%XMM15 |
(238) 0x43c8ea VMOVQ %XMM15,%RSI |
(238) 0x43c8ef SUB %R14,%RSI |
(238) 0x43c8f2 VMOVUPD (%RDI,%RSI,8),%ZMM15 |
(238) 0x43c8f9 VMOVUPD 0x8(%RDI,%RSI,8),%ZMM16 |
(238) 0x43c904 VMOVUPD (%R8,%RSI,8),%ZMM17 |
(238) 0x43c90b VMOVUPD 0x8(%R8,%RSI,8),%ZMM18 |
(238) 0x43c916 VMOVUPD (%R12,%RSI,8),%ZMM19 |
(238) 0x43c91d VMOVUPD 0x8(%R12,%RSI,8),%ZMM20 |
(238) 0x43c928 VMOVUPD (%R11,%RSI,8),%ZMM21 |
(238) 0x43c92f VMOVUPD (%R9,%RSI,8),%ZMM22 |
(238) 0x43c936 VMULPD (%RAX,%RSI,8),%ZMM22,%ZMM23 |
(238) 0x43c93d VFMADD231PD (%R15,%RSI,8),%ZMM23,%ZMM11 |
(238) 0x43c944 VMOVUPD 0x8(%R11,%RSI,8),%ZMM24 |
(238) 0x43c94f VFMADD231PD (%RDX,%RSI,8),%ZMM22,%ZMM13 |
(238) 0x43c956 VMULPD %ZMM15,%ZMM15,%ZMM15 |
(238) 0x43c95c VFMADD213PD %ZMM15,%ZMM17,%ZMM17 |
(238) 0x43c962 VFMADD231PD %ZMM16,%ZMM16,%ZMM17 |
(238) 0x43c968 VFMADD231PD %ZMM18,%ZMM18,%ZMM17 |
(238) 0x43c96e VFMADD213PD %ZMM17,%ZMM19,%ZMM19 |
(238) 0x43c974 VFMADD213PD %ZMM19,%ZMM21,%ZMM21 |
(238) 0x43c97a VFMADD231PD %ZMM20,%ZMM20,%ZMM21 |
(238) 0x43c980 VFMADD231PD %ZMM24,%ZMM24,%ZMM21 |
(238) 0x43c986 VADDPD %ZMM22,%ZMM9,%ZMM9 |
(238) 0x43c98c VADDPD %ZMM23,%ZMM10,%ZMM10 |
(238) 0x43c992 VMULPD %ZMM21,%ZMM23,%ZMM15 |
(238) 0x43c998 VFMADD231PD %ZMM7,%ZMM15,%ZMM12 |
(238) 0x43c99e VPADDQ %ZMM6,%ZMM14,%ZMM14 |
(238) 0x43c9a4 ADD $0x8,%R10 |
(238) 0x43c9a8 CMP %RBX,%R10 |
(238) 0x43c9ab JB 43c8e0 |
(237) 0x43c9b1 VEXTRACTF64X4 $0x1,%ZMM13,%YMM14 |
(237) 0x43c9b8 VADDPD %ZMM14,%ZMM13,%ZMM13 |
(237) 0x43c9be VMOVAPD %XMM13,%XMM14 |
(237) 0x43c9c3 VEXTRACTF128 $0x1,%YMM13,%XMM13 |
(237) 0x43c9c9 VADDPD %XMM13,%XMM14,%XMM13 |
(237) 0x43c9ce VSHUFPD $0x1,%XMM13,%XMM13,%XMM14 |
(237) 0x43c9d4 VADDSD %XMM14,%XMM13,%XMM13 |
(237) 0x43c9d9 VADDSD %XMM3,%XMM13,%XMM3 |
(237) 0x43c9dd VEXTRACTF64X4 $0x1,%ZMM12,%YMM13 |
(237) 0x43c9e4 VADDPD %ZMM13,%ZMM12,%ZMM12 |
(237) 0x43c9ea VMOVAPD %XMM12,%XMM13 |
(237) 0x43c9ef VEXTRACTF128 $0x1,%YMM12,%XMM12 |
(237) 0x43c9f5 VADDPD %XMM12,%XMM13,%XMM12 |
(237) 0x43c9fa VSHUFPD $0x1,%XMM12,%XMM12,%XMM13 |
(237) 0x43ca00 VADDSD %XMM13,%XMM12,%XMM12 |
(237) 0x43ca05 VADDSD %XMM0,%XMM12,%XMM0 |
(237) 0x43ca09 VEXTRACTF64X4 $0x1,%ZMM11,%YMM12 |
(237) 0x43ca10 VADDPD %ZMM12,%ZMM11,%ZMM11 |
(237) 0x43ca16 VMOVAPD %XMM11,%XMM12 |
(237) 0x43ca1b VEXTRACTF128 $0x1,%YMM11,%XMM11 |
(237) 0x43ca21 VADDPD %XMM11,%XMM12,%XMM11 |
(237) 0x43ca26 VSHUFPD $0x1,%XMM11,%XMM11,%XMM12 |
(237) 0x43ca2c VADDSD %XMM12,%XMM11,%XMM11 |
(237) 0x43ca31 VADDSD %XMM2,%XMM11,%XMM2 |
(237) 0x43ca35 VEXTRACTF64X4 $0x1,%ZMM10,%YMM11 |
(237) 0x43ca3c VADDPD %ZMM11,%ZMM10,%ZMM10 |
(237) 0x43ca42 VMOVAPD %XMM10,%XMM11 |
(237) 0x43ca47 VEXTRACTF128 $0x1,%YMM10,%XMM10 |
(237) 0x43ca4d VADDPD %XMM10,%XMM11,%XMM10 |
(237) 0x43ca52 VSHUFPD $0x1,%XMM10,%XMM10,%XMM11 |
(237) 0x43ca58 VADDSD %XMM11,%XMM10,%XMM10 |
(237) 0x43ca5d VADDSD %XMM4,%XMM10,%XMM4 |
(237) 0x43ca61 VEXTRACTF64X4 $0x1,%ZMM9,%YMM10 |
(237) 0x43ca68 VADDPD %ZMM10,%ZMM9,%ZMM9 |
(237) 0x43ca6e VMOVAPD %XMM9,%XMM10 |
(237) 0x43ca73 VEXTRACTF128 $0x1,%YMM9,%XMM9 |
(237) 0x43ca79 VADDPD %XMM9,%XMM10,%XMM9 |
(237) 0x43ca7e VSHUFPD $0x1,%XMM9,%XMM9,%XMM10 |
(237) 0x43ca84 VADDSD %XMM10,%XMM9,%XMM9 |
(237) 0x43ca89 VADDSD %XMM1,%XMM9,%XMM1 |
(237) 0x43ca8d MOV %RBX,%RDX |
(237) 0x43ca90 CMP -0x48(%RBP),%RBX |
(237) 0x43ca94 MOV -0x50(%RBP),%R9D |
(237) 0x43ca98 MOV 0xa0(%RBP),%RDI |
(237) 0x43ca9f MOV 0x98(%RBP),%RSI |
(237) 0x43caa6 MOV 0x90(%RBP),%R15 |
(237) 0x43caad MOV 0x88(%RBP),%R12 |
(237) 0x43cab4 MOV %RBX,%R8 |
(237) 0x43cab7 MOV -0xb8(%RBP),%R11 |
(237) 0x43cabe MOV -0xb0(%RBP),%RBX |
(237) 0x43cac5 JE 43c7bf |
(237) 0x43cacb MOV -0x48(%RBP),%RAX |
(237) 0x43cacf SUB %RDX,%RAX |
(237) 0x43cad2 VPBROADCASTQ %RAX,%ZMM10 |
(237) 0x43cad8 MOV -0xc8(%RBP),%R10 |
(237) 0x43cadf SUB -0x98(%RBP),%R10 |
(237) 0x43cae6 MOV %R13,%RDI |
(237) 0x43cae9 MOV %RCX,%R8 |
(237) 0x43caec LEA 0x1(%R10),%RAX |
(237) 0x43caf0 IMUL %RAX,%R13 |
(237) 0x43caf4 IMUL %RAX,%RCX |
(237) 0x43caf8 IMUL %R10,%RDI |
(237) 0x43cafc IMUL %R10,%R8 |
(237) 0x43cb00 IMUL %R10,%RBX |
(237) 0x43cb04 MOV -0xd8(%RBP),%R12 |
(237) 0x43cb0b IMUL %R10,%R12 |
(237) 0x43cb0f MOV -0xa0(%RBP),%RSI |
(237) 0x43cb16 IMUL %R10,%RSI |
(237) 0x43cb1a MOV -0xd0(%RBP),%R15 |
(237) 0x43cb21 IMUL %R10,%R15 |
(237) 0x43cb25 MOV 0x30(%RBP),%RAX |
(237) 0x43cb29 ADD %RAX,%RDI |
(237) 0x43cb2c MOV %RBX,%R10 |
(237) 0x43cb2f MOV 0x28(%RBP),%RBX |
(237) 0x43cb33 ADD %RBX,%R8 |
(237) 0x43cb36 ADD %RAX,%R13 |
(237) 0x43cb39 ADD %RBX,%RCX |
(237) 0x43cb3c ADD 0x50(%RBP),%R10 |
(237) 0x43cb40 ADD 0x48(%RBP),%R12 |
(237) 0x43cb44 ADD 0x40(%RBP),%RSI |
(237) 0x43cb48 ADD 0x38(%RBP),%R15 |
(237) 0x43cb4c VXORPD %XMM14,%XMM14,%XMM14 |
(237) 0x43cb51 VXORPD %XMM13,%XMM13,%XMM13 |
(237) 0x43cb56 VXORPD %XMM12,%XMM12,%XMM12 |
(237) 0x43cb5b VXORPD %XMM11,%XMM11,%XMM11 |
(237) 0x43cb60 VXORPD %XMM9,%XMM9,%XMM9 |
(237) 0x43cb65 VMOVDQA64 %ZMM5,%ZMM15 |
(237) 0x43cb6b JMP 43cbb5 |
0x43cb6d NOPW %CS:(%RAX,%RAX,1) |
0x43cb7c NOPL (%RAX) |
(239) 0x43cb80 VMOVAPD %ZMM16,%ZMM9{%K1} |
(239) 0x43cb86 VMOVAPD %ZMM17,%ZMM11{%K1} |
(239) 0x43cb8c VMOVAPD %ZMM18,%ZMM12{%K1} |
(239) 0x43cb92 VMOVAPD %ZMM19,%ZMM13{%K1} |
(239) 0x43cb98 VMOVAPD %ZMM20,%ZMM14{%K1} |
(239) 0x43cb9e VPADDQ %ZMM6,%ZMM15,%ZMM15 |
(239) 0x43cba4 VPCMPLTUQ %ZMM10,%ZMM15,%K0 |
(239) 0x43cbab KORTESTB %K0,%K0 |
(239) 0x43cbaf JE 43c6c0 |
(239) 0x43cbb5 VPCMPLTUQ %ZMM10,%ZMM15,%K1 |
(239) 0x43cbbc KORTESTB %K1,%K1 |
(239) 0x43cbc0 VXORPD %XMM16,%XMM16,%XMM16 |
(239) 0x43cbc6 VXORPD %XMM17,%XMM17,%XMM17 |
(239) 0x43cbcc VXORPD %XMM18,%XMM18,%XMM18 |
(239) 0x43cbd2 VXORPD %XMM19,%XMM19,%XMM19 |
(239) 0x43cbd8 VXORPD %XMM20,%XMM20,%XMM20 |
(239) 0x43cbde JE 43cb80 |
(239) 0x43cbe0 VMOVDQA64 %XMM15,%XMM16 |
(239) 0x43cbe6 VMOVQ %XMM15,%RAX |
(239) 0x43cbeb ADD %RDX,%RAX |
(239) 0x43cbee ADD -0xe0(%RBP),%RAX |
(239) 0x43cbf5 SUB %R14,%RAX |
(239) 0x43cbf8 VMOVUPD (%RDI,%RAX,8),%ZMM16{%K1}{z} |
(239) 0x43cbff VMULPD %ZMM16,%ZMM16,%ZMM16 |
(239) 0x43cc05 VMOVUPD (%R8,%RAX,8),%ZMM17{%K1}{z} |
(239) 0x43cc0c VMOVUPD 0x8(%RDI,%RAX,8),%ZMM18{%K1}{z} |
(239) 0x43cc17 VMOVUPD 0x8(%R8,%RAX,8),%ZMM19{%K1}{z} |
(239) 0x43cc22 VFMADD213PD %ZMM16,%ZMM17,%ZMM17 |
(239) 0x43cc28 VFMADD213PD %ZMM17,%ZMM18,%ZMM18 |
(239) 0x43cc2e VMOVUPD (%R13,%RAX,8),%ZMM16{%K1}{z} |
(239) 0x43cc36 VMOVUPD (%RCX,%RAX,8),%ZMM17{%K1}{z} |
(239) 0x43cc3d VFMADD231PD %ZMM19,%ZMM19,%ZMM18 |
(239) 0x43cc43 VFMADD213PD %ZMM18,%ZMM16,%ZMM16 |
(239) 0x43cc49 VMOVUPD 0x8(%R13,%RAX,8),%ZMM19{%K1}{z} |
(239) 0x43cc54 VMOVUPD 0x8(%RCX,%RAX,8),%ZMM18{%K1}{z} |
(239) 0x43cc5f VFMADD231PD %ZMM17,%ZMM17,%ZMM16 |
(239) 0x43cc65 VFMADD213PD %ZMM16,%ZMM19,%ZMM19 |
(239) 0x43cc6b VMOVUPD (%R10,%RAX,8),%ZMM21{%K1}{z} |
(239) 0x43cc72 VMOVUPD (%R12,%RAX,8),%ZMM16{%K1}{z} |
(239) 0x43cc79 VFMADD231PD %ZMM18,%ZMM18,%ZMM19 |
(239) 0x43cc7f VMULPD %ZMM21,%ZMM16,%ZMM20 |
(239) 0x43cc85 VADDPD %ZMM21,%ZMM9,%ZMM16 |
(239) 0x43cc8b VADDPD %ZMM20,%ZMM11,%ZMM17 |
(239) 0x43cc91 VMOVUPD (%RSI,%RAX,8),%ZMM18{%K1}{z} |
(239) 0x43cc98 VFMADD213PD %ZMM12,%ZMM20,%ZMM18 |
(239) 0x43cc9e VMULPD %ZMM19,%ZMM20,%ZMM19 |
(239) 0x43cca4 VFMADD132PD %ZMM7,%ZMM13,%ZMM19 |
(239) 0x43ccaa VMOVUPD (%R15,%RAX,8),%ZMM20{%K1}{z} |
(239) 0x43ccb1 VFMADD213PD %ZMM14,%ZMM21,%ZMM20 |
(239) 0x43ccb7 JMP 43cb80 |
0x43ccbc NOPL (%RAX) |
0x43ccc0 MOV -0x30(%RBP),%RDI |
0x43ccc4 JNE 43cd61 |
0x43ccca VMOVSD -0x80(%RBP),%XMM0 |
0x43cccf MOV (%RDI),%ESI |
0x43ccd1 MOV $0x74a5f0,%EDI |
0x43ccd6 MOV %R15,%RDX |
0x43ccd9 CALL 404290 <__kmpc_atomic_float8_add@plt> |
0x43ccde VMOVSD -0x78(%RBP),%XMM0 |
0x43cce3 MOV -0x30(%RBP),%RAX |
0x43cce7 MOV (%RAX),%ESI |
0x43cce9 MOV $0x74a610,%EDI |
0x43ccee MOV %R14,%RDX |
0x43ccf1 CALL 404290 <__kmpc_atomic_float8_add@plt> |
0x43ccf6 VMOVSD -0x70(%RBP),%XMM0 |
0x43ccfb MOV -0x30(%RBP),%RAX |
0x43ccff MOV (%RAX),%ESI |
0x43cd01 MOV $0x74a630,%EDI |
0x43cd06 MOV -0x88(%RBP),%RDX |
0x43cd0d CALL 404290 <__kmpc_atomic_float8_add@plt> |
0x43cd12 VMOVSD -0x68(%RBP),%XMM0 |
0x43cd17 MOV -0x30(%RBP),%RAX |
0x43cd1b MOV (%RAX),%ESI |
0x43cd1d MOV $0x74a650,%EDI |
0x43cd22 MOV %R13,%RDX |
0x43cd25 CALL 404290 <__kmpc_atomic_float8_add@plt> |
0x43cd2a VMOVSD -0x60(%RBP),%XMM0 |
0x43cd2f MOV -0x30(%RBP),%RAX |
0x43cd33 MOV (%RAX),%ESI |
0x43cd35 MOV $0x74a670,%EDI |
0x43cd3a MOV -0x90(%RBP),%RDX |
0x43cd41 CALL 404290 <__kmpc_atomic_float8_add@plt> |
0x43cd46 MOV -0x30(%RBP),%RAX |
0x43cd4a MOV (%RAX),%ESI |
0x43cd4c MOV $0x7a3060,%RDX |
0x43cd53 MOV $0x74a6d0,%EDI |
0x43cd58 CALL 404900 <__kmpc_end_reduce@plt> |
0x43cd5d MOV -0x30(%RBP),%RDI |
0x43cd61 MOV (%RDI),%ESI |
0x43cd63 MOV $0x74a6f0,%EDI |
0x43cd68 CALL 404660 <__kmpc_barrier@plt> |
0x43cd6d ADD $0xb8,%RSP |
0x43cd74 POP %RBX |
0x43cd75 POP %R12 |
0x43cd77 POP %R13 |
0x43cd79 POP %R14 |
0x43cd7b POP %R15 |
0x43cd7d POP %RBP |
0x43cd7e RET |
0x43cd7f NOP |
Path / |
Source file and lines | field_summary_kernel.f90:54-74 |
Module | exec |
nb instructions | 191 |
nb uops | 202 |
loop length | 894 |
used x86 registers | 15 |
used mmx registers | 0 |
used xmm registers | 6 |
used ymm registers | 1 |
used zmm registers | 4 |
nb stack references | 31 |
micro-operation queue | 33.67 cycles |
front end | 33.67 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 9.70 | 9.70 | 21.67 | 21.67 | 23.00 | 9.50 | 9.50 | 23.00 | 23.00 | 23.00 | 9.60 | 21.67 |
cycles | 9.70 | 9.70 | 21.67 | 21.67 | 23.00 | 9.50 | 9.50 | 23.00 | 23.00 | 23.00 | 9.60 | 21.67 |
Cycles executing div or sqrt instructions | NA |
FE+BE cycles | 32.00 |
Stall cycles | 0.00 |
Front-end | 33.67 |
Dispatch | 23.00 |
Overall L1 | 33.67 |
all | 8% |
load | 4% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 19% |
all | 16% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 83% |
all | 11% |
load | 2% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 33% |
all | 11% |
load | 12% |
store | 10% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 6% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 12% |
all | 14% |
load | 12% |
store | 12% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 22% |
all | 12% |
load | 12% |
store | 11% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 10% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 14% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
SUB $0xb8,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R9,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R8,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x60(%RBP),%EBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x58(%RBP),%EAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
SUB %EBX,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOVL $0,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JS 43cd61 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x931> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %RCX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %RDX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %RDI,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOVL $0,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %EAX,-0x34(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOVL $0x1,-0x54(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
SUB $0x8,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
LEA -0x54(%RBP),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x58(%RBP),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x38(%RBP),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x34(%RBP),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV $0x74a5b0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %ESI,-0x4c(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV $0x22,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
PUSH $0x1 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH $0x1 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RAX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
CALL 4045a0 <__kmpc_for_static_init_4@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
ADD $0x20,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x38(%RBP),%EAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x34(%RBP),%R9D | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VXORPD %XMM0,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VXORPD %XMM2,%XMM2,%XMM2 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VXORPD %XMM3,%XMM3,%XMM3 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VXORPD %XMM4,%XMM4,%XMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VXORPD %XMM1,%XMM1,%XMM1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
SUB %EAX,%R9D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 43c600 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x1d0> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x20(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x18(%RBP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x10(%RBP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD %XMM1,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD %XMM4,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD %XMM3,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD %XMM2,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD %XMM0,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV $0x74a5d0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x4c(%RBP),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 404190 <__kmpc_for_static_fini@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
SUB $0x8,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV $0x7a3060,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA -0x80(%RBP),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV $0x74a690,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x43c3e0,%R9D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x5,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x28,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
PUSH %RAX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
CALL 404780 <__kmpc_reduce@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
ADD $0x10,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP $0x2,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 43ccc0 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x890> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP $0x1,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x30(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JNE 43cd61 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x931> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VMOVSD -0x80(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%R15),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%R15) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD -0x78(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%R14),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%R14) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD -0x70(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x88(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%RAX),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%RAX) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD -0x68(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%R13),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%R13) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD -0x60(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x90(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%RAX),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%RAX) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x7a3060,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x74a6b0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JMP 43cd58 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x928> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0xa0(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x98(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x90(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x88(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x70(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x68(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD $-0x2,%R14D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOVSXD %R14D,%R14 | 1 | 0 | 0.33 | 0 | 0 | 0 | 0.33 | 0 | 0 | 0 | 0 | 0.33 | 0 | 1 | 0.33 |
ADD $-0x2,%R13D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOVSXD %R13D,%RDX | 1 | 0 | 0.33 | 0 | 0 | 0 | 0.33 | 0 | 0 | 0 | 0 | 0.33 | 0 | 1 | 0.33 |
MOV %RDX,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x40(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %EBX,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RDX,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOVSXD (%RAX),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RCX),%EAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
SUB %EDX,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RAX,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
INC %EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMP $0x2,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMOVGE %EAX,%ECX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %RDX,-0xe0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VPBROADCASTD %EDX,%YMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV %RCX,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %ECX,%R8D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $0x7ffffff8,%R8D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
VPMOVSXDQ %YMM0,%ZMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVDQA64 0xcf777(%RIP),%ZMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.50 |
VPBROADCASTQ 0xcfae5(%RIP),%ZMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.33 |
VBROADCASTSD 0xcfad3(%RIP),%ZMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.33 |
VMOVDQA %XMM0,%XMM8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VPXOR %XMM0,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %R11D,%R11D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R9D,-0x50(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R8,-0xa8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JMP 43c7cf <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x39f> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x30(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JNE 43cd61 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x931> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VMOVSD -0x80(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a5f0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R15,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
VMOVSD -0x78(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a610,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R14,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
VMOVSD -0x70(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a630,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x88(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
VMOVSD -0x68(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a650,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R13,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
VMOVSD -0x60(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a670,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x90(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x7a3060,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x74a6d0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 404900 <__kmpc_end_reduce@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a6f0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 404660 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
ADD $0xb8,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
RET | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 0 | 2.13 |
NOP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
Source file and lines | field_summary_kernel.f90:54-74 |
Module | exec |
nb instructions | 191 |
nb uops | 202 |
loop length | 894 |
used x86 registers | 15 |
used mmx registers | 0 |
used xmm registers | 6 |
used ymm registers | 1 |
used zmm registers | 4 |
nb stack references | 31 |
micro-operation queue | 33.67 cycles |
front end | 33.67 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 9.70 | 9.70 | 21.67 | 21.67 | 23.00 | 9.50 | 9.50 | 23.00 | 23.00 | 23.00 | 9.60 | 21.67 |
cycles | 9.70 | 9.70 | 21.67 | 21.67 | 23.00 | 9.50 | 9.50 | 23.00 | 23.00 | 23.00 | 9.60 | 21.67 |
Cycles executing div or sqrt instructions | NA |
FE+BE cycles | 32.00 |
Stall cycles | 0.00 |
Front-end | 33.67 |
Dispatch | 23.00 |
Overall L1 | 33.67 |
all | 8% |
load | 4% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 19% |
all | 16% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 83% |
all | 11% |
load | 2% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 33% |
all | 11% |
load | 12% |
store | 10% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 6% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 12% |
all | 14% |
load | 12% |
store | 12% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 22% |
all | 12% |
load | 12% |
store | 11% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 10% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 14% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
SUB $0xb8,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R9,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R8,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x60(%RBP),%EBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x58(%RBP),%EAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
SUB %EBX,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOVL $0,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JS 43cd61 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x931> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %RCX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %RDX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %RDI,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOVL $0,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %EAX,-0x34(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOVL $0x1,-0x54(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
SUB $0x8,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
LEA -0x54(%RBP),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x58(%RBP),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x38(%RBP),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x34(%RBP),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV $0x74a5b0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %ESI,-0x4c(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV $0x22,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
PUSH $0x1 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH $0x1 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RAX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
CALL 4045a0 <__kmpc_for_static_init_4@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
ADD $0x20,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x38(%RBP),%EAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x34(%RBP),%R9D | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VXORPD %XMM0,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VXORPD %XMM2,%XMM2,%XMM2 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VXORPD %XMM3,%XMM3,%XMM3 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VXORPD %XMM4,%XMM4,%XMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VXORPD %XMM1,%XMM1,%XMM1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
SUB %EAX,%R9D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 43c600 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x1d0> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x20(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x18(%RBP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x10(%RBP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD %XMM1,-0x80(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD %XMM4,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD %XMM3,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD %XMM2,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD %XMM0,-0x60(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV $0x74a5d0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x4c(%RBP),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
CALL 404190 <__kmpc_for_static_fini@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
SUB $0x8,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV $0x7a3060,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA -0x80(%RBP),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV $0x74a690,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x43c3e0,%R9D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x5,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x28,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
PUSH %RAX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
CALL 404780 <__kmpc_reduce@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
ADD $0x10,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP $0x2,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 43ccc0 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x890> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP $0x1,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x30(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JNE 43cd61 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x931> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VMOVSD -0x80(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%R15),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%R15) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD -0x78(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%R14),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%R14) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD -0x70(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x88(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%RAX),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%RAX) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD -0x68(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%R13),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%R13) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD -0x60(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x90(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VADDSD (%RAX),%XMM0,%XMM0 | 1 | 0 | 0.50 | 0.33 | 0.33 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.50 |
VMOVSD %XMM0,(%RAX) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x7a3060,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x74a6b0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JMP 43cd58 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x928> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0xa0(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x98(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x90(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x88(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x70(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x68(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD $-0x2,%R14D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOVSXD %R14D,%R14 | 1 | 0 | 0.33 | 0 | 0 | 0 | 0.33 | 0 | 0 | 0 | 0 | 0.33 | 0 | 1 | 0.33 |
ADD $-0x2,%R13D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOVSXD %R13D,%RDX | 1 | 0 | 0.33 | 0 | 0 | 0 | 0.33 | 0 | 0 | 0 | 0 | 0.33 | 0 | 1 | 0.33 |
MOV %RDX,-0x98(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x40(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %EBX,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RDX,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOVSXD (%RAX),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RCX),%EAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
SUB %EDX,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %RAX,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
INC %EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMP $0x2,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMOVGE %EAX,%ECX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
MOV %RDX,-0xe0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VPBROADCASTD %EDX,%YMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV %RCX,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %ECX,%R8D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $0x7ffffff8,%R8D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
VPMOVSXDQ %YMM0,%ZMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVDQA64 0xcf777(%RIP),%ZMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.50 |
VPBROADCASTQ 0xcfae5(%RIP),%ZMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.33 |
VBROADCASTSD 0xcfad3(%RIP),%ZMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.33 |
VMOVDQA %XMM0,%XMM8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VPXOR %XMM0,%XMM0,%XMM0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %R11D,%R11D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R9D,-0x50(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R8,-0xa8(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JMP 43c7cf <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x39f> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x30(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JNE 43cd61 <field_summary_kernel_module_mp_field_summary_kernel_.DIR.OMP.PARALLEL.2+0x931> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VMOVSD -0x80(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a5f0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R15,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
VMOVSD -0x78(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a610,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R14,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
VMOVSD -0x70(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a630,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x88(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
VMOVSD -0x68(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a650,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R13,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
VMOVSD -0x60(%RBP),%XMM0 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a670,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x90(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CALL 404290 <__kmpc_atomic_float8_add@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x7a3060,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x74a6d0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 404900 <__kmpc_end_reduce@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RDI),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74a6f0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 404660 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
ADD $0xb8,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
RET | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 0 | 2.13 |
NOP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
Name | Coverage (%) | Time (s) |
---|---|---|
▼field_summary_kernel_.DIR.OMP.PARALLEL.2– | 0.29 | 0.22 |
▼Loop 237 - field_summary_kernel.f90:56-71 - exec– | 0 | 0 |
○Loop 238 - field_summary_kernel.f90:58-71 - exec | 0.29 | 0.22 |
○Loop 239 - field_summary_kernel.f90:58-71 - exec | 0 | 0 |