Function: ideal_gas_kernel._omp_fn.0 | Module: exec | Source: ideal_gas_kernel.f90:45-55 | Coverage: 4.92% |
---|
Function: ideal_gas_kernel._omp_fn.0 | Module: exec | Source: ideal_gas_kernel.f90:45-55 | Coverage: 4.92% |
---|
/home/eoseret/qaas_runs_CPU_9468/171-152-3172/intel/CloverLeafFC/build/CloverLeafFC/CloverLeaf_ref/kernels/ideal_gas_kernel.f90: 45 - 55 |
-------------------------------------------------------------------------------- |
45: !$OMP PARALLEL |
46: !$OMP DO PRIVATE(v,pressurebyenergy,pressurebyvolume,sound_speed_squared) |
47: DO k=y_min,y_max |
48: !$OMP SIMD |
49: DO j=x_min,x_max |
50: v=1.0_8/density(j,k) |
51: pressure(j,k)=(1.4_8-1.0_8)*density(j,k)*energy(j,k) |
52: pressurebyenergy=(1.4_8-1.0_8)*density(j,k) |
53: pressurebyvolume=-density(j,k)*pressure(j,k) |
54: sound_speed_squared=v*v*(pressure(j,k)*pressurebyenergy-pressurebyvolume) |
55: soundspeed(j,k)=SQRT(sound_speed_squared) |
0x442750 PUSH %RBP |
0x442751 MOV %RSP,%RBP |
0x442754 PUSH %R15 |
0x442756 PUSH %R14 |
0x442758 PUSH %R13 |
0x44275a PUSH %R12 |
0x44275c PUSH %RBX |
0x44275d AND $-0x40,%RSP |
0x442761 ADD $-0x80,%RSP |
0x442765 MOV 0x70(%RDI),%RDX |
0x442769 MOV 0x68(%RDI),%RCX |
0x44276d MOV %RDI,0x78(%RSP) |
0x442772 MOV 0x78(%RDI),%RAX |
0x442776 MOV 0x60(%RDI),%RBX |
0x44277a MOV 0x50(%RDI),%RSI |
0x44277e MOV 0x40(%RDI),%R8 |
0x442782 MOV %RDX,0x48(%RSP) |
0x442787 MOV 0x10(%RDI),%R9 |
0x44278b MOV 0x58(%RDI),%R13 |
0x44278f MOV %RCX,0x60(%RSP) |
0x442794 MOV 0x48(%RDI),%R14 |
0x442798 MOV %RSI,0x38(%RSP) |
0x44279d MOV (%R9),%R12D |
0x4427a0 MOV %R8,0x30(%RSP) |
0x4427a5 MOV %RAX,0x70(%RSP) |
0x4427aa MOV %RBX,0x40(%RSP) |
0x4427af CALL 402080 <@plt_start@+0x60> |
0x4427b4 MOV %EAX,%R15D |
0x4427b7 CALL 402180 <@plt_start@+0x160> |
0x4427bc MOV 0x78(%RSP),%RCX |
0x4427c1 MOV %EAX,%EDI |
0x4427c3 MOV 0x18(%RCX),%R10 |
0x4427c7 MOV (%R10),%EAX |
0x4427ca INC %EAX |
0x4427cc SUB %R12D,%EAX |
0x4427cf CLTD |
0x4427d0 IDIV %R15D |
0x4427d3 CMP %EDX,%EDI |
0x4427d5 JL 442df3 |
0x4427db IMUL %EAX,%EDI |
0x4427de ADD %EDX,%EDI |
0x4427e0 ADD %EDI,%EAX |
0x4427e2 CMP %EAX,%EDI |
0x4427e4 JGE 442dc5 |
0x4427ea MOV 0x8(%RCX),%R8 |
0x4427ee MOV 0x30(%RCX),%R9 |
0x4427f2 ADD %R12D,%EAX |
0x4427f5 ADD %R12D,%EDI |
0x4427f8 MOV (%RCX),%R11 |
0x4427fb MOV 0x28(%RCX),%R10 |
0x4427ff MOV %EAX,0x2c(%RSP) |
0x442803 MOV (%R8),%R15D |
0x442806 VMOVSD 0x57552(%RIP),%XMM7 |
0x44280e MOV %R9,0x78(%RSP) |
0x442813 MOV 0x30(%RSP),%R9 |
0x442818 MOVSXD (%R11),%RSI |
0x44281b MOV %EDI,0x6c(%RSP) |
0x44281f LEA 0x1(%R15),%EAX |
0x442823 MOV 0x40(%RSP),%R11 |
0x442828 VMOVSD 0x57a10(%RIP),%XMM6 |
0x442830 MOV %R15D,0x68(%RSP) |
0x442835 MOV %EAX,0x24(%RSP) |
0x442839 MOVSXD %EDI,%RAX |
0x44283c ADD %RSI,%R14 |
0x44283f MOV 0x60(%RSP),%RDI |
0x442844 IMUL %RAX,%R9 |
0x442848 LEA (%R13,%RSI,1),%R8 |
0x44284d MOV 0x38(%RSP),%R13 |
0x442852 MOV %RSI,%RBX |
0x442855 IMUL %RAX,%R11 |
0x442859 ADD %RSI,%RDI |
0x44285c MOV %ESI,0x58(%RSP) |
0x442860 VMOVSD 0x579e0(%RIP),%XMM5 |
0x442868 IMUL %RAX,%R13 |
0x44286c MOV %RSI,0x10(%RSP) |
0x442871 MOV 0x70(%RSP),%RSI |
0x442876 VBROADCASTSD %XMM7,%YMM10 |
0x44287b ADD %R14,%R9 |
0x44287e MOV 0x48(%RSP),%R14 |
0x442883 MOV %R10,0x70(%RSP) |
0x442888 VBROADCASTSD %XMM6,%YMM9 |
0x44288d ADD %RBX,%RSI |
0x442890 ADD %R11,%RDI |
0x442893 MOV 0x20(%RCX),%RDX |
0x442897 MOV 0x38(%RCX),%R12 |
0x44289b IMUL %R14,%RAX |
0x44289f VBROADCASTSD %XMM5,%YMM8 |
0x4428a4 VBROADCASTSD %XMM7,%ZMM4 |
0x4428aa ADD %R13,%R8 |
0x4428ad VBROADCASTSD %XMM6,%ZMM3 |
0x4428b3 VBROADCASTSD %XMM5,%ZMM2 |
0x4428b9 ADD %RAX,%RSI |
0x4428bc MOV %R15D,%EAX |
0x4428bf SUB %EBX,%EAX |
0x4428c1 MOV %EAX,0x5c(%RSP) |
0x4428c5 INC %EAX |
0x4428c7 MOV %EAX,%R13D |
0x4428ca MOV %EAX,%R11D |
0x4428cd SHR $0x3,%R13D |
0x4428d1 AND $-0x8,%R11D |
0x4428d5 SAL $0x6,%R13 |
0x4428d9 CMP %R15D,%EBX |
0x4428dc LEA (%R11,%RBX,1),%R14D |
0x4428e0 MOV %R11D,0x20(%RSP) |
0x4428e5 CMOVLE 0x24(%RSP),%EBX |
0x4428ea AND $0x7,%EAX |
0x4428ed MOV %R13,0x60(%RSP) |
0x4428f2 XOR %R11D,%R11D |
0x4428f5 MOV %R14D,0x1c(%RSP) |
0x4428fa MOV %EBX,0x18(%RSP) |
0x4428fe MOV %EAX,0x28(%RSP) |
0x442902 MOV %RCX,0x8(%RSP) |
0x442907 NOPW (%RAX,%RAX,1) |
(224) 0x442910 MOV 0x68(%RSP),%ECX |
(224) 0x442914 CMP %ECX,0x58(%RSP) |
(224) 0x442918 JG 442dd8 |
(224) 0x44291e CMPL $0x6,0x5c(%RSP) |
(224) 0x442923 JBE 442de8 |
(224) 0x442929 MOV 0x70(%RSP),%R10 |
(224) 0x44292e MOV 0x60(%RSP),%R14 |
(224) 0x442933 LEA (%RDX,%R9,8),%R15 |
(224) 0x442937 LEA (%R12,%RSI,8),%RCX |
(224) 0x44293b MOV 0x78(%RSP),%R11 |
(224) 0x442940 XOR %EAX,%EAX |
(224) 0x442942 LEA (%R10,%R8,8),%R13 |
(224) 0x442946 LEA -0x40(%R14),%R10 |
(224) 0x44294a SHR $0x6,%R10 |
(224) 0x44294e LEA (%R11,%RDI,8),%RBX |
(224) 0x442952 INC %R10 |
(224) 0x442955 AND $0x3,%R10D |
(224) 0x442959 JE 442a5a |
(224) 0x44295f CMP $0x1,%R10 |
(224) 0x442963 JE 442a03 |
(224) 0x442969 CMP $0x2,%R10 |
(224) 0x44296d JE 4429b7 |
(224) 0x44296f VMOVUPD (%R15),%ZMM1 |
(224) 0x442975 MOV $0x40,%EAX |
(224) 0x44297a VMULPD (%R13),%ZMM1,%ZMM0 |
(224) 0x442981 VDIVPD %ZMM1,%ZMM4,%ZMM11 |
(224) 0x442987 VMULPD %ZMM11,%ZMM11,%ZMM14 |
(224) 0x44298d VMULPD %ZMM3,%ZMM0,%ZMM12 |
(224) 0x442993 VMOVUPD %ZMM12,(%RBX) |
(224) 0x442999 VMULPD (%R15),%ZMM2,%ZMM13 |
(224) 0x44299f VMULPD %ZMM14,%ZMM13,%ZMM15 |
(224) 0x4429a5 VMULPD %ZMM12,%ZMM15,%ZMM1 |
(224) 0x4429ab VSQRTPD %ZMM1,%ZMM11 |
(224) 0x4429b1 VMOVUPD %ZMM11,(%RCX) |
(224) 0x4429b7 VMOVUPD (%R15,%RAX,1),%ZMM0 |
(224) 0x4429be VMULPD (%R13,%RAX,1),%ZMM0,%ZMM13 |
(224) 0x4429c6 VDIVPD %ZMM0,%ZMM4,%ZMM12 |
(224) 0x4429cc VMULPD %ZMM12,%ZMM12,%ZMM1 |
(224) 0x4429d2 VMULPD %ZMM3,%ZMM13,%ZMM14 |
(224) 0x4429d8 VMOVUPD %ZMM14,(%RBX,%RAX,1) |
(224) 0x4429df VMULPD (%R15,%RAX,1),%ZMM2,%ZMM15 |
(224) 0x4429e6 VMULPD %ZMM1,%ZMM15,%ZMM11 |
(224) 0x4429ec VMULPD %ZMM14,%ZMM11,%ZMM0 |
(224) 0x4429f2 VSQRTPD %ZMM0,%ZMM12 |
(224) 0x4429f8 VMOVUPD %ZMM12,(%RCX,%RAX,1) |
(224) 0x4429ff ADD $0x40,%RAX |
(224) 0x442a03 VMOVUPD (%R15,%RAX,1),%ZMM13 |
(224) 0x442a0a VMULPD (%R13,%RAX,1),%ZMM13,%ZMM15 |
(224) 0x442a12 VDIVPD %ZMM13,%ZMM4,%ZMM14 |
(224) 0x442a18 VMULPD %ZMM14,%ZMM14,%ZMM11 |
(224) 0x442a1e VMULPD %ZMM3,%ZMM15,%ZMM1 |
(224) 0x442a24 VMOVUPD %ZMM1,(%RBX,%RAX,1) |
(224) 0x442a2b VMULPD (%R15,%RAX,1),%ZMM2,%ZMM0 |
(224) 0x442a32 VMULPD %ZMM11,%ZMM0,%ZMM12 |
(224) 0x442a38 VMULPD %ZMM1,%ZMM12,%ZMM13 |
(224) 0x442a3e VSQRTPD %ZMM13,%ZMM14 |
(224) 0x442a44 VMOVUPD %ZMM14,(%RCX,%RAX,1) |
(224) 0x442a4b ADD $0x40,%RAX |
(224) 0x442a4f CMP %RAX,0x60(%RSP) |
(224) 0x442a54 JE 442b97 |
(225) 0x442a5a VMOVUPD (%R15,%RAX,1),%ZMM15 |
(225) 0x442a61 VMULPD (%R13,%RAX,1),%ZMM15,%ZMM1 |
(225) 0x442a69 VDIVPD %ZMM15,%ZMM4,%ZMM11 |
(225) 0x442a6f VMULPD %ZMM11,%ZMM11,%ZMM13 |
(225) 0x442a75 VMULPD %ZMM3,%ZMM1,%ZMM12 |
(225) 0x442a7b VMOVUPD %ZMM12,(%RBX,%RAX,1) |
(225) 0x442a82 VMULPD (%R15,%RAX,1),%ZMM2,%ZMM0 |
(225) 0x442a89 VMULPD %ZMM13,%ZMM0,%ZMM14 |
(225) 0x442a8f VMULPD %ZMM12,%ZMM14,%ZMM15 |
(225) 0x442a95 VSQRTPD %ZMM15,%ZMM11 |
(225) 0x442a9b VMOVUPD %ZMM11,(%RCX,%RAX,1) |
(225) 0x442aa2 VMOVUPD 0x40(%R15,%RAX,1),%ZMM1 |
(225) 0x442aaa VMULPD 0x40(%R13,%RAX,1),%ZMM1,%ZMM0 |
(225) 0x442ab2 VDIVPD %ZMM1,%ZMM4,%ZMM12 |
(225) 0x442ab8 VMULPD %ZMM12,%ZMM12,%ZMM15 |
(225) 0x442abe VMULPD %ZMM3,%ZMM0,%ZMM13 |
(225) 0x442ac4 VMOVUPD %ZMM13,0x40(%RBX,%RAX,1) |
(225) 0x442acc VMULPD 0x40(%R15,%RAX,1),%ZMM2,%ZMM14 |
(225) 0x442ad4 VMULPD %ZMM15,%ZMM14,%ZMM11 |
(225) 0x442ada VMULPD %ZMM13,%ZMM11,%ZMM1 |
(225) 0x442ae0 VSQRTPD %ZMM1,%ZMM12 |
(225) 0x442ae6 VMOVUPD %ZMM12,0x40(%RCX,%RAX,1) |
(225) 0x442aee VMOVUPD 0x80(%R15,%RAX,1),%ZMM0 |
(225) 0x442af6 VMULPD 0x80(%R13,%RAX,1),%ZMM0,%ZMM14 |
(225) 0x442afe VDIVPD %ZMM0,%ZMM4,%ZMM13 |
(225) 0x442b04 VMULPD %ZMM13,%ZMM13,%ZMM11 |
(225) 0x442b0a VMULPD %ZMM3,%ZMM14,%ZMM15 |
(225) 0x442b10 VMOVUPD %ZMM15,0x80(%RBX,%RAX,1) |
(225) 0x442b18 VMULPD 0x80(%R15,%RAX,1),%ZMM2,%ZMM1 |
(225) 0x442b20 VMULPD %ZMM11,%ZMM1,%ZMM12 |
(225) 0x442b26 VMULPD %ZMM15,%ZMM12,%ZMM0 |
(225) 0x442b2c VSQRTPD %ZMM0,%ZMM13 |
(225) 0x442b32 VMOVUPD %ZMM13,0x80(%RCX,%RAX,1) |
(225) 0x442b3a VMOVUPD 0xc0(%R15,%RAX,1),%ZMM14 |
(225) 0x442b42 VMULPD 0xc0(%R13,%RAX,1),%ZMM14,%ZMM1 |
(225) 0x442b4a VDIVPD %ZMM14,%ZMM4,%ZMM15 |
(225) 0x442b50 VMULPD %ZMM15,%ZMM15,%ZMM11 |
(225) 0x442b56 VMULPD %ZMM3,%ZMM1,%ZMM12 |
(225) 0x442b5c VMOVUPD %ZMM12,0xc0(%RBX,%RAX,1) |
(225) 0x442b64 VMULPD 0xc0(%R15,%RAX,1),%ZMM2,%ZMM0 |
(225) 0x442b6c VMULPD %ZMM11,%ZMM0,%ZMM13 |
(225) 0x442b72 VMULPD %ZMM12,%ZMM13,%ZMM14 |
(225) 0x442b78 VSQRTPD %ZMM14,%ZMM15 |
(225) 0x442b7e VMOVUPD %ZMM15,0xc0(%RCX,%RAX,1) |
(225) 0x442b86 ADD $0x100,%RAX |
(225) 0x442b8c CMP %RAX,0x60(%RSP) |
(225) 0x442b91 JNE 442a5a |
(224) 0x442b97 MOV 0x28(%RSP),%R15D |
(224) 0x442b9c TEST %R15D,%R15D |
(224) 0x442b9f JE 442d66 |
(224) 0x442ba5 MOV 0x20(%RSP),%ECX |
(224) 0x442ba9 MOV 0x1c(%RSP),%EAX |
(224) 0x442bad MOV 0x5c(%RSP),%R13D |
(224) 0x442bb2 SUB %ECX,%R13D |
(224) 0x442bb5 LEA 0x1(%R13),%R10D |
(224) 0x442bb9 CMP $0x2,%R13D |
(224) 0x442bbd JBE 442c26 |
(224) 0x442bbf LEA (%R9,%RCX,1),%R11 |
(224) 0x442bc3 MOV 0x70(%RSP),%R14 |
(224) 0x442bc8 LEA (%R8,%RCX,1),%R15 |
(224) 0x442bcc MOV 0x78(%RSP),%R13 |
(224) 0x442bd1 LEA (%RDX,%R11,8),%R11 |
(224) 0x442bd5 LEA (%RCX,%RSI,1),%RBX |
(224) 0x442bd9 ADD %RDI,%RCX |
(224) 0x442bdc VMOVUPD (%R11),%YMM1 |
(224) 0x442be1 VDIVPD %YMM1,%YMM10,%YMM12 |
(224) 0x442be5 VMULPD (%R14,%R15,8),%YMM1,%YMM0 |
(224) 0x442beb VMULPD %YMM9,%YMM0,%YMM13 |
(224) 0x442bf0 VMOVUPD %YMM13,(%R13,%RCX,8) |
(224) 0x442bf7 VMULPD (%R11),%YMM8,%YMM14 |
(224) 0x442bfc VMULPD %YMM12,%YMM12,%YMM11 |
(224) 0x442c01 VMULPD %YMM11,%YMM14,%YMM15 |
(224) 0x442c06 VMULPD %YMM13,%YMM15,%YMM1 |
(224) 0x442c0b VSQRTPD %YMM1,%YMM12 |
(224) 0x442c0f VMOVUPD %YMM12,(%R12,%RBX,8) |
(224) 0x442c15 TEST $0x3,%R10B |
(224) 0x442c19 JE 442d66 |
(224) 0x442c1f AND $-0x4,%R10D |
(224) 0x442c23 ADD %R10D,%EAX |
(224) 0x442c26 MOV 0x10(%RSP),%RCX |
(224) 0x442c2b MOV %R9,%R10 |
(224) 0x442c2e MOV %R8,%R13 |
(224) 0x442c31 MOV %RDI,%RBX |
(224) 0x442c34 MOV %RSI,%R11 |
(224) 0x442c37 SUB %RCX,%R11 |
(224) 0x442c3a SUB %RCX,%R10 |
(224) 0x442c3d SUB %RCX,%R13 |
(224) 0x442c40 SUB %RCX,%RBX |
(224) 0x442c43 MOVSXD %EAX,%RCX |
(224) 0x442c46 MOV %R11,0x50(%RSP) |
(224) 0x442c4b MOV 0x70(%RSP),%R11 |
(224) 0x442c50 LEA (%R10,%RCX,1),%R14 |
(224) 0x442c54 LEA (%RCX,%R13,1),%R15 |
(224) 0x442c58 VMOVSD (%RDX,%R14,8),%XMM0 |
(224) 0x442c5e VDIVSD %XMM0,%XMM7,%XMM13 |
(224) 0x442c62 VMULSD (%R11,%R15,8),%XMM0,%XMM14 |
(224) 0x442c68 MOV 0x78(%RSP),%R11 |
(224) 0x442c6d LEA (%RCX,%RBX,1),%R15 |
(224) 0x442c71 VMULSD %XMM6,%XMM14,%XMM11 |
(224) 0x442c75 VMOVSD %XMM11,(%R11,%R15,8) |
(224) 0x442c7b VMULSD (%RDX,%R14,8),%XMM5,%XMM15 |
(224) 0x442c81 MOV 0x50(%RSP),%R14 |
(224) 0x442c86 ADD %R14,%RCX |
(224) 0x442c89 VMULSD %XMM11,%XMM15,%XMM1 |
(224) 0x442c8e VMULSD %XMM13,%XMM13,%XMM12 |
(224) 0x442c93 VMULSD %XMM12,%XMM1,%XMM0 |
(224) 0x442c98 VSQRTSD %XMM0,%XMM0,%XMM0 |
(224) 0x442c9c VMOVSD %XMM0,(%R12,%RCX,8) |
(224) 0x442ca2 LEA 0x1(%RAX),%ECX |
(224) 0x442ca5 CMP %ECX,0x68(%RSP) |
(224) 0x442ca9 JL 442d66 |
(224) 0x442caf MOVSXD %ECX,%RCX |
(224) 0x442cb2 MOV 0x70(%RSP),%R11 |
(224) 0x442cb7 ADD $0x2,%EAX |
(224) 0x442cba LEA (%RCX,%R10,1),%R14 |
(224) 0x442cbe LEA (%RCX,%R13,1),%R15 |
(224) 0x442cc2 VMOVSD (%RDX,%R14,8),%XMM13 |
(224) 0x442cc8 VDIVSD %XMM13,%XMM7,%XMM14 |
(224) 0x442ccd VMULSD (%R11,%R15,8),%XMM13,%XMM11 |
(224) 0x442cd3 MOV 0x78(%RSP),%R11 |
(224) 0x442cd8 LEA (%RCX,%RBX,1),%R15 |
(224) 0x442cdc VMULSD %XMM6,%XMM11,%XMM15 |
(224) 0x442ce0 VMOVSD %XMM15,(%R11,%R15,8) |
(224) 0x442ce6 VMULSD (%RDX,%R14,8),%XMM5,%XMM1 |
(224) 0x442cec MOV 0x50(%RSP),%R14 |
(224) 0x442cf1 ADD %R14,%RCX |
(224) 0x442cf4 VMULSD %XMM15,%XMM1,%XMM12 |
(224) 0x442cf9 VMULSD %XMM14,%XMM14,%XMM0 |
(224) 0x442cfe VMULSD %XMM0,%XMM12,%XMM13 |
(224) 0x442d02 VSQRTSD %XMM13,%XMM13,%XMM13 |
(224) 0x442d07 VMOVSD %XMM13,(%R12,%RCX,8) |
(224) 0x442d0d CMP %EAX,0x68(%RSP) |
(224) 0x442d11 JL 442d66 |
(224) 0x442d13 CLTQ |
(224) 0x442d15 MOV 0x70(%RSP),%RCX |
(224) 0x442d1a ADD %RAX,%R10 |
(224) 0x442d1d ADD %RAX,%R13 |
(224) 0x442d20 ADD %RAX,%RBX |
(224) 0x442d23 ADD %RAX,%R14 |
(224) 0x442d26 VMOVSD (%RDX,%R10,8),%XMM14 |
(224) 0x442d2c VDIVSD %XMM14,%XMM7,%XMM11 |
(224) 0x442d31 VMULSD (%RCX,%R13,8),%XMM14,%XMM15 |
(224) 0x442d37 MOV 0x78(%RSP),%R13 |
(224) 0x442d3c VMULSD %XMM6,%XMM15,%XMM1 |
(224) 0x442d40 VMOVSD %XMM1,(%R13,%RBX,8) |
(224) 0x442d47 VMULSD (%RDX,%R10,8),%XMM5,%XMM12 |
(224) 0x442d4d VMULSD %XMM1,%XMM12,%XMM0 |
(224) 0x442d51 VMULSD %XMM11,%XMM11,%XMM13 |
(224) 0x442d56 VMULSD %XMM13,%XMM0,%XMM14 |
(224) 0x442d5b VSQRTSD %XMM14,%XMM14,%XMM14 |
(224) 0x442d60 VMOVSD %XMM14,(%R12,%R14,8) |
(224) 0x442d66 MOV 0x18(%RSP),%EAX |
(224) 0x442d6a MOV $0x1,%R11D |
(224) 0x442d70 MOV %EAX,0x50(%RSP) |
(224) 0x442d74 NOPL (%RAX) |
(224) 0x442d78 INCL 0x6c(%RSP) |
(224) 0x442d7c MOV 0x30(%RSP),%RBX |
(224) 0x442d81 MOV 0x38(%RSP),%R14 |
(224) 0x442d86 MOV 0x40(%RSP),%RCX |
(224) 0x442d8b MOV 0x48(%RSP),%R13 |
(224) 0x442d90 ADD %RBX,%R9 |
(224) 0x442d93 ADD %R14,%R8 |
(224) 0x442d96 ADD %RCX,%RDI |
(224) 0x442d99 MOV 0x6c(%RSP),%R10D |
(224) 0x442d9e ADD %R13,%RSI |
(224) 0x442da1 CMP %R10D,0x2c(%RSP) |
(224) 0x442da6 JG 442910 |
0x442dac MOV 0x8(%RSP),%RDX |
0x442db1 TEST %R11B,%R11B |
0x442db4 JE 442dfc |
0x442db6 MOV 0x50(%RSP),%R12D |
0x442dbb MOV %R12D,0x80(%RDX) |
0x442dc2 VZEROUPPER |
0x442dc5 LEA -0x28(%RBP),%RSP |
0x442dc9 POP %RBX |
0x442dca POP %R12 |
0x442dcc POP %R13 |
0x442dce POP %R14 |
0x442dd0 POP %R15 |
0x442dd2 POP %RBP |
0x442dd3 RET |
0x442dd4 NOPL (%RAX) |
(224) 0x442dd8 MOV 0x18(%RSP),%EBX |
(224) 0x442ddc CMP %EBX,0x24(%RSP) |
(224) 0x442de0 JNE 442d78 |
(224) 0x442de2 JMP 442d66 |
0x442de4 NOPL (%RAX) |
(224) 0x442de8 MOV 0x58(%RSP),%EAX |
(224) 0x442dec XOR %ECX,%ECX |
(224) 0x442dee JMP 442bad |
0x442df3 INC %EAX |
0x442df5 XOR %EDX,%EDX |
0x442df7 JMP 4427db |
0x442dfc VZEROUPPER |
0x442dff LEA -0x28(%RBP),%RSP |
0x442e03 POP %RBX |
0x442e04 POP %R12 |
0x442e06 POP %R13 |
0x442e08 POP %R14 |
0x442e0a POP %R15 |
0x442e0c POP %RBP |
0x442e0d RET |
0x442e0e XCHG %AX,%AX |
Path / |
Source file and lines | ideal_gas_kernel.f90:45-55 |
Module | exec |
nb instructions | 142 |
nb uops | 149 |
loop length | 525 |
used x86 registers | 16 |
used mmx registers | 0 |
used xmm registers | 3 |
used ymm registers | 3 |
used zmm registers | 3 |
nb stack references | 21 |
micro-operation queue | 24.83 cycles |
front end | 24.83 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 9.90 | 9.80 | 15.67 | 15.67 | 15.50 | 9.80 | 9.70 | 15.50 | 15.50 | 15.50 | 9.80 | 15.67 |
cycles | 9.90 | 13.33 | 15.67 | 15.67 | 15.50 | 9.80 | 9.70 | 15.50 | 15.50 | 15.50 | 9.80 | 15.67 |
Cycles executing div or sqrt instructions | 6.00 |
FE+BE cycles | 24.25-25.63 |
Stall cycles | 0.00-1.24 |
Front-end | 24.83 |
Dispatch | 15.67 |
DIV/SQRT | 6.00 |
Overall L1 | 24.83 |
all | 5% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 28% |
all | 0% |
load | 0% |
store | NA (no store vectorizable/vectorized instructions) |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 0% |
all | 4% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 0% |
other | 16% |
all | 10% |
load | 10% |
store | 9% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 11% |
all | 12% |
load | 12% |
store | NA (no store vectorizable/vectorized instructions) |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 12% |
all | 10% |
load | 10% |
store | 9% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 6% |
other | 12% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
AND $-0x40,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
ADD $-0x80,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV 0x70(%RDI),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x68(%RDI),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RDI,0x78(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x78(%RDI),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x60(%RDI),%RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x50(%RDI),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x40(%RDI),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RDX,0x48(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x10(%RDI),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x58(%RDI),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x60(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x48(%RDI),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RSI,0x38(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV (%R9),%R12D | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R8,0x30(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RAX,0x70(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RBX,0x40(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CALL 402080 <@plt_start@+0x60> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %EAX,%R15D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 402180 <@plt_start@+0x160> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0x78(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %EAX,%EDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV 0x18(%RCX),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%R10),%EAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
INC %EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SUB %R12D,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CLTD | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
IDIV %R15D | 4 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 6 |
CMP %EDX,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JL 442df3 <__ideal_gas_kernel_module_MOD_ideal_gas_kernel._omp_fn.0+0x6a3> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
IMUL %EAX,%EDI | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %EDX,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %EDI,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMP %EAX,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 442dc5 <__ideal_gas_kernel_module_MOD_ideal_gas_kernel._omp_fn.0+0x675> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x8(%RCX),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x30(%RCX),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %R12D,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %R12D,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV (%RCX),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x28(%RCX),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %EAX,0x2c(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV (%R8),%R15D | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD 0x57552(%RIP),%XMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R9,0x78(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x30(%RSP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOVSXD (%R11),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %EDI,0x6c(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
LEA 0x1(%R15),%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV 0x40(%RSP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD 0x57a10(%RIP),%XMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R15D,0x68(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %EAX,0x24(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOVSXD %EDI,%RAX | 1 | 0 | 0.33 | 0 | 0 | 0 | 0.33 | 0 | 0 | 0 | 0 | 0.33 | 0 | 1 | 0.33 |
ADD %RSI,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x60(%RSP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %RAX,%R9 | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA (%R13,%RSI,1),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x38(%RSP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RSI,%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
IMUL %RAX,%R11 | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %RSI,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %ESI,0x58(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD 0x579e0(%RIP),%XMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %RAX,%R13 | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV %RSI,0x10(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x70(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VBROADCASTSD %XMM7,%YMM10 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %R14,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x48(%RSP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R10,0x70(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VBROADCASTSD %XMM6,%YMM9 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %RBX,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %R11,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x20(%RCX),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%RCX),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %R14,%RAX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM5,%YMM8 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM7,%ZMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %R13,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VBROADCASTSD %XMM6,%ZMM3 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM5,%ZMM2 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %RAX,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R15D,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %EBX,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %EAX,0x5c(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
INC %EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %EAX,%R13D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %EAX,%R11D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SHR $0x3,%R13D | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
AND $-0x8,%R11D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
SAL $0x6,%R13 | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
CMP %R15D,%EBX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA (%R11,%RBX,1),%R14D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV %R11D,0x20(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CMOVLE 0x24(%RSP),%EBX | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.50 |
AND $0x7,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV %R13,0x60(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
XOR %R11D,%R11D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R14D,0x1c(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %EBX,0x18(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %EAX,0x28(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,0x8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x8(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R11B,%R11B | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JE 442dfc <__ideal_gas_kernel_module_MOD_ideal_gas_kernel._omp_fn.0+0x6ac> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x50(%RSP),%R12D | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R12D,0x80(%RDX) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
LEA -0x28(%RBP),%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
RET | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 0 | 2.13 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
INC %EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 4427db <__ideal_gas_kernel_module_MOD_ideal_gas_kernel._omp_fn.0+0x8b> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
LEA -0x28(%RBP),%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
RET | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 0 | 2.13 |
XCHG %AX,%AX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
Source file and lines | ideal_gas_kernel.f90:45-55 |
Module | exec |
nb instructions | 142 |
nb uops | 149 |
loop length | 525 |
used x86 registers | 16 |
used mmx registers | 0 |
used xmm registers | 3 |
used ymm registers | 3 |
used zmm registers | 3 |
nb stack references | 21 |
micro-operation queue | 24.83 cycles |
front end | 24.83 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 9.90 | 9.80 | 15.67 | 15.67 | 15.50 | 9.80 | 9.70 | 15.50 | 15.50 | 15.50 | 9.80 | 15.67 |
cycles | 9.90 | 13.33 | 15.67 | 15.67 | 15.50 | 9.80 | 9.70 | 15.50 | 15.50 | 15.50 | 9.80 | 15.67 |
Cycles executing div or sqrt instructions | 6.00 |
FE+BE cycles | 24.25-25.63 |
Stall cycles | 0.00-1.24 |
Front-end | 24.83 |
Dispatch | 15.67 |
DIV/SQRT | 6.00 |
Overall L1 | 24.83 |
all | 5% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 28% |
all | 0% |
load | 0% |
store | NA (no store vectorizable/vectorized instructions) |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 0% |
all | 4% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 0% |
other | 16% |
all | 10% |
load | 10% |
store | 9% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 11% |
all | 12% |
load | 12% |
store | NA (no store vectorizable/vectorized instructions) |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 12% |
all | 10% |
load | 10% |
store | 9% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 6% |
other | 12% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
AND $-0x40,%RSP | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
ADD $-0x80,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV 0x70(%RDI),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x68(%RDI),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RDI,0x78(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x78(%RDI),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x60(%RDI),%RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x50(%RDI),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x40(%RDI),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RDX,0x48(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x10(%RDI),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x58(%RDI),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x60(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x48(%RDI),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RSI,0x38(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV (%R9),%R12D | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R8,0x30(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RAX,0x70(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RBX,0x40(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CALL 402080 <@plt_start@+0x60> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %EAX,%R15D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 402180 <@plt_start@+0x160> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0x78(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %EAX,%EDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV 0x18(%RCX),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%R10),%EAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
INC %EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
SUB %R12D,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CLTD | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
IDIV %R15D | 4 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 6 |
CMP %EDX,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JL 442df3 <__ideal_gas_kernel_module_MOD_ideal_gas_kernel._omp_fn.0+0x6a3> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
IMUL %EAX,%EDI | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %EDX,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %EDI,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CMP %EAX,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 442dc5 <__ideal_gas_kernel_module_MOD_ideal_gas_kernel._omp_fn.0+0x675> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x8(%RCX),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x30(%RCX),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %R12D,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %R12D,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV (%RCX),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x28(%RCX),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %EAX,0x2c(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV (%R8),%R15D | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD 0x57552(%RIP),%XMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R9,0x78(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x30(%RSP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOVSXD (%R11),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %EDI,0x6c(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
LEA 0x1(%R15),%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV 0x40(%RSP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD 0x57a10(%RIP),%XMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R15D,0x68(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %EAX,0x24(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOVSXD %EDI,%RAX | 1 | 0 | 0.33 | 0 | 0 | 0 | 0.33 | 0 | 0 | 0 | 0 | 0.33 | 0 | 1 | 0.33 |
ADD %RSI,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x60(%RSP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %RAX,%R9 | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA (%R13,%RSI,1),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x38(%RSP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RSI,%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
IMUL %RAX,%R11 | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %RSI,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %ESI,0x58(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMOVSD 0x579e0(%RIP),%XMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %RAX,%R13 | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV %RSI,0x10(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x70(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VBROADCASTSD %XMM7,%YMM10 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %R14,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x48(%RSP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R10,0x70(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VBROADCASTSD %XMM6,%YMM9 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %RBX,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %R11,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x20(%RCX),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%RCX),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %R14,%RAX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM5,%YMM8 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM7,%ZMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %R13,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VBROADCASTSD %XMM6,%ZMM3 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM5,%ZMM2 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %RAX,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R15D,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %EBX,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %EAX,0x5c(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
INC %EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %EAX,%R13D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %EAX,%R11D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SHR $0x3,%R13D | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
AND $-0x8,%R11D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
SAL $0x6,%R13 | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
CMP %R15D,%EBX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA (%R11,%RBX,1),%R14D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV %R11D,0x20(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CMOVLE 0x24(%RSP),%EBX | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.50 |
AND $0x7,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
MOV %R13,0x60(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
XOR %R11D,%R11D | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R14D,0x1c(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %EBX,0x18(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %EAX,0x28(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,0x8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x8(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R11B,%R11B | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JE 442dfc <__ideal_gas_kernel_module_MOD_ideal_gas_kernel._omp_fn.0+0x6ac> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x50(%RSP),%R12D | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R12D,0x80(%RDX) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
LEA -0x28(%RBP),%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
RET | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 0 | 2.13 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
INC %EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 4427db <__ideal_gas_kernel_module_MOD_ideal_gas_kernel._omp_fn.0+0x8b> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
VZEROUPPER | 2 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
LEA -0x28(%RBP),%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
RET | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 0 | 2.13 |
XCHG %AX,%AX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
Name | Coverage (%) | Time (s) |
---|---|---|
▼ideal_gas_kernel._omp_fn.0– | 4.92 | 3.62 |
▼Loop 224 - ideal_gas_kernel.f90:45-55 - exec– | 0 | 0 |
○Loop 225 - ideal_gas_kernel.f90:50-55 - exec | 4.92 | 3.62 |