Loop Id: 774 | Module: exec | Source: MultiBsplineRef.hpp:227-262 [...] | Coverage: 0.04% |
---|
Loop Id: 774 | Module: exec | Source: MultiBsplineRef.hpp:227-262 [...] | Coverage: 0.04% |
---|
0x436180 MOV 0x30(%RSP),%RCX |
0x436185 VMOVUPD 0x620(%RSP),%YMM8 |
0x43618e VMOVUPD 0x600(%RSP),%YMM14 |
0x436197 VMOVUPD 0x5e0(%RSP),%YMM15 |
0x4361a0 VMOVUPD 0x5c0(%RSP),%YMM30 |
0x4361a8 VMOVUPD 0x5a0(%RSP),%YMM29 |
0x4361b0 VMOVUPD 0x580(%RSP),%YMM11 |
0x4361b9 VMOVUPD 0x380(%RSP),%YMM7 |
0x4361c2 VMOVUPD 0x360(%RSP),%YMM5 |
0x4361cb VMOVSD 0x180(%RSP),%XMM28 |
0x4361d3 VMOVSD 0x178(%RSP),%XMM24 |
0x4361db MOV 0x90(%RSP),%R10 |
0x4361e3 MOV 0x18(%RSP),%R8 |
0x4361e8 MOV 0x1a8(%RSP),%RDI |
0x4361f0 LEA 0x1(%RDI),%RAX |
0x4361f4 MOV 0x308(%RSP),%RDX |
0x4361fc ADD %RDX,0xe0(%RSP) |
0x436204 ADD %RDX,0xe8(%RSP) |
0x43620c ADD %RDX,0xf0(%RSP) |
0x436214 ADD %RDX,0xf8(%RSP) |
0x43621c MOV 0xc8(%RSP),%RDX |
0x436224 ADD %RDX,0x188(%RSP) |
0x43622c ADD %RDX,0x190(%RSP) |
0x436234 ADD %RDX,0x198(%RSP) |
0x43623c ADD %RDX,0x1a0(%RSP) |
0x436244 CMP $0x3,%RDI |
0x436248 JE 436030 |
0x43624e VMOVSD 0x540(%RSP,%RAX,8),%XMM16 |
0x436259 VMULSD %XMM28,%XMM16,%XMM27 |
0x43625f VMULSD %XMM24,%XMM16,%XMM25 |
0x436265 VMOVSD 0x310(%RSP),%XMM4 |
0x43626e VMULSD %XMM4,%XMM16,%XMM16 |
0x436274 VMOVSD 0x3e0(%RSP,%RAX,8),%XMM17 |
0x43627c VMULSD %XMM24,%XMM17,%XMM26 |
0x436282 VMULSD %XMM4,%XMM17,%XMM17 |
0x436288 MOV %RAX,0x1a8(%RSP) |
0x436290 VMULSD 0x3a0(%RSP,%RAX,8),%XMM4,%XMM4 |
0x436299 VMOVUPD %XMM4,0x260(%RSP) |
0x4362a2 CMP $0x81,%R10D |
0x4362a9 MOV 0xd8(%RSP),%RAX |
0x4362b1 MOV 0xd0(%RSP),%RDX |
0x4362b9 JB 436d10 |
0x4362bf TEST %R10D,%R10D |
0x4362c2 JLE 4361e8 |
0x4362c8 VBROADCASTSD %XMM27,%YMM19 |
0x4362ce VBROADCASTSD %XMM26,%YMM23 |
0x4362d4 VBROADCASTSD %XMM25,%YMM24 |
0x4362da VMOVUPD 0x260(%RSP),%XMM11 |
0x4362e3 VBROADCASTSD %XMM11,%YMM21 |
0x4362e9 VBROADCASTSD %XMM17,%YMM22 |
0x4362ef VBROADCASTSD %XMM16,%YMM28 |
0x4362f5 MOV 0x80(%RSP),%RCX |
0x4362fd MOV %RCX,0x258(%RSP) |
0x436305 MOV 0x78(%RSP),%RCX |
0x43630a MOV %RCX,0x250(%RSP) |
0x436312 MOV 0x68(%RSP),%RCX |
0x436317 MOV %RCX,0x248(%RSP) |
0x43631f MOV 0x70(%RSP),%RCX |
0x436324 MOV %RCX,0x240(%RSP) |
0x43632c MOV 0x60(%RSP),%RCX |
0x436331 MOV %RCX,0x238(%RSP) |
0x436339 MOV %R8,%RDI |
0x43633c MOV %R13,0x210(%RSP) |
0x436344 MOV %RDX,0x208(%RSP) |
0x43634c MOV 0x170(%RSP),%RCX |
0x436354 MOV %RCX,0x200(%RSP) |
0x43635c MOV 0xc0(%RSP),%RCX |
0x436364 MOV %RCX,0x1f8(%RSP) |
0x43636c MOV 0x158(%RSP),%RCX |
0x436374 MOV %RCX,0x1f0(%RSP) |
0x43637c MOV 0x160(%RSP),%RCX |
0x436384 MOV %RCX,0x1e8(%RSP) |
0x43638c MOV 0x168(%RSP),%RCX |
0x436394 MOV %RCX,0x1e0(%RSP) |
0x43639c MOV %R9,0x108(%RSP) |
0x4363a4 MOV 0x1a0(%RSP),%RCX |
0x4363ac MOV %RCX,0x230(%RSP) |
0x4363b4 MOV 0x198(%RSP),%RCX |
0x4363bc MOV %RCX,0x228(%RSP) |
0x4363c4 MOV 0x190(%RSP),%RCX |
0x4363cc MOV %RCX,0x220(%RSP) |
0x4363d4 MOV 0x188(%RSP),%RCX |
0x4363dc MOV %RCX,0x218(%RSP) |
0x4363e4 MOV 0x300(%RSP),%R15 |
0x4363ec MOV 0xf8(%RSP),%RCX |
0x4363f4 MOV %RCX,0x1d8(%RSP) |
0x4363fc MOV 0xf0(%RSP),%RCX |
0x436404 MOV %RCX,0x1d0(%RSP) |
0x43640c MOV 0xe8(%RSP),%RCX |
0x436414 MOV %RCX,0x1c8(%RSP) |
0x43641c MOV 0xe0(%RSP),%RCX |
0x436424 MOV %RCX,0x1c0(%RSP) |
0x43642c MOV %RAX,0x100(%RSP) |
0x436434 XOR %EBX,%EBX |
0x436436 VMOVUPD 0x600(%RSP),%YMM13 |
0x43643f VMOVUPD 0x5e0(%RSP),%YMM14 |
0x436448 VMOVUPD 0x5c0(%RSP),%YMM15 |
0x436451 JMP 4365b6 |
(778) 0x436460 ADDQ $0x200,0x100(%RSP) |
(778) 0x43646c ADDQ $0x200,0x1c0(%RSP) |
(778) 0x436478 ADDQ $0x200,0x1c8(%RSP) |
(778) 0x436484 ADDQ $0x200,0x1d0(%RSP) |
(778) 0x436490 ADDQ $0x200,0x1d8(%RSP) |
(778) 0x43649c MOV 0x330(%RSP),%R15 |
(778) 0x4364a4 ADD $-0x40,%R15 |
(778) 0x4364a8 ADDQ $0x40,0x218(%RSP) |
(778) 0x4364b1 ADDQ $0x40,0x220(%RSP) |
(778) 0x4364ba ADDQ $0x40,0x228(%RSP) |
(778) 0x4364c3 ADDQ $0x40,0x230(%RSP) |
(778) 0x4364cc ADDQ $0x200,0x108(%RSP) |
(778) 0x4364d8 ADDQ $0x200,0x1e0(%RSP) |
(778) 0x4364e4 ADDQ $0x200,0x1e8(%RSP) |
(778) 0x4364f0 ADDQ $0x200,0x1f0(%RSP) |
(778) 0x4364fc ADDQ $0x200,0x1f8(%RSP) |
(778) 0x436508 ADDQ $0x200,0x200(%RSP) |
(778) 0x436514 ADDQ $0x200,0x208(%RSP) |
(778) 0x436520 ADDQ $0x200,0x210(%RSP) |
(778) 0x43652c MOV 0x110(%RSP),%RDI |
(778) 0x436534 ADD $0x200,%RDI |
(778) 0x43653b ADDQ $0x40,0x238(%RSP) |
(778) 0x436544 ADDQ $0x40,0x240(%RSP) |
(778) 0x43654d ADDQ $0x40,0x248(%RSP) |
(778) 0x436556 ADDQ $0x40,0x250(%RSP) |
(778) 0x43655f ADDQ $0x40,0x258(%RSP) |
(778) 0x436568 MOV 0x328(%RSP),%RAX |
(778) 0x436570 CMP 0x318(%RSP),%RAX |
(778) 0x436578 LEA 0x1(%RAX),%RBX |
(778) 0x43657c MOV 0x28(%RSP),%R9 |
(778) 0x436581 MOV 0x88(%RSP),%R11 |
(778) 0x436589 VMOVUPD 0x420(%RSP),%YMM10 |
(778) 0x436592 VMOVUPD 0x340(%RSP),%YMM9 |
(778) 0x43659b MOV 0x20(%RSP),%RSI |
(778) 0x4365a0 MOV 0x1b8(%RSP),%R13 |
(778) 0x4365a8 MOV 0x320(%RSP),%R12 |
(778) 0x4365b0 JE 436180 |
(778) 0x4365b6 MOV %RDI,0x110(%RSP) |
(778) 0x4365be CMP $0x3f,%R15 |
(778) 0x4365c2 MOV $0x3f,%ECX |
(778) 0x4365c7 CMOVL %R15,%RCX |
(778) 0x4365cb MOV %RBX,%R11 |
(778) 0x4365ce SAL $0x6,%R11 |
(778) 0x4365d2 NOT %R11 |
(778) 0x4365d5 ADD 0x90(%RSP),%R11 |
(778) 0x4365dd CMP $0x3f,%R11 |
(778) 0x4365e1 MOV $0x3f,%EAX |
(778) 0x4365e6 CMOVGE %RAX,%R11 |
(778) 0x4365ea INC %R11 |
(778) 0x4365ed MOV %R11,%R10 |
(778) 0x4365f0 AND $-0x4,%R10 |
(778) 0x4365f4 MOV %R15,0x330(%RSP) |
(778) 0x4365fc MOV %RBX,0x328(%RSP) |
(778) 0x436604 MOV %RCX,0x120(%RSP) |
(778) 0x43660c MOV %R10,0x38(%RSP) |
(778) 0x436611 MOV %R11,0x118(%RSP) |
(778) 0x436619 JE 4367f0 |
(778) 0x43661f LEA -0x1(%R10),%RAX |
(778) 0x436623 XOR %ECX,%ECX |
(778) 0x436625 VMOVUPD 0x620(%RSP),%YMM11 |
(778) 0x43662e VMOVUPD 0x5a0(%RSP),%YMM0 |
(778) 0x436637 VMOVUPD 0x580(%RSP),%YMM1 |
(778) 0x436640 VMOVUPD 0x380(%RSP),%YMM2 |
(778) 0x436649 VMOVUPD 0x360(%RSP),%YMM3 |
(778) 0x436652 VMOVAPD %YMM31,%YMM9 |
(778) 0x436658 VMOVUPD 0x460(%RSP),%YMM31 |
(778) 0x436660 VMOVUPD 0x440(%RSP),%YMM18 |
(778) 0x436668 MOV 0x1d8(%RSP),%RDX |
(778) 0x436670 MOV 0x1d0(%RSP),%RSI |
(778) 0x436678 MOV 0x1c8(%RSP),%RDI |
(778) 0x436680 MOV 0x1c0(%RSP),%R8 |
(778) 0x436688 MOV 0x100(%RSP),%R9 |
(780) 0x436690 VMOVUPD (%RDX,%RCX,8),%YMM29 |
(780) 0x436697 VMOVUPD (%RSI,%RCX,8),%YMM30 |
(780) 0x43669e VMOVUPD (%RDI,%RCX,8),%YMM4 |
(780) 0x4366a3 VMOVUPD (%R8,%RCX,8),%YMM5 |
(780) 0x4366a9 VMULPD %YMM11,%YMM29,%YMM6 |
(780) 0x4366af VFMADD231PD %YMM13,%YMM30,%YMM6 |
(780) 0x4366b5 VFMADD231PD %YMM14,%YMM4,%YMM6 |
(780) 0x4366ba VFMADD231PD %YMM5,%YMM15,%YMM6 |
(780) 0x4366bf VMULPD %YMM0,%YMM29,%YMM7 |
(780) 0x4366c5 VMULPD %YMM1,%YMM30,%YMM8 |
(780) 0x4366cb VFMADD231PD %YMM3,%YMM5,%YMM8 |
(780) 0x4366d0 VFMADD231PD %YMM2,%YMM4,%YMM7 |
(780) 0x4366d5 VFMADD231PD %YMM8,%YMM12,%YMM7 |
(780) 0x4366da VMULPD %YMM31,%YMM29,%YMM8 |
(780) 0x4366e0 VFMADD231PD %YMM30,%YMM9,%YMM8 |
(780) 0x4366e6 VFMADD231PD %YMM4,%YMM18,%YMM8 |
(780) 0x4366ec VFMADD231PD %YMM5,%YMM12,%YMM8 |
(780) 0x4366f1 VMULPD %YMM6,%YMM19,%YMM4 |
(780) 0x4366f7 VMOVUPD %YMM4,0x1948(%RSP,%RCX,8) |
(780) 0x436700 VMULPD %YMM6,%YMM23,%YMM4 |
(780) 0x436706 VMOVUPD %YMM4,0x1748(%RSP,%RCX,8) |
(780) 0x43670f VMULPD %YMM24,%YMM7,%YMM4 |
(780) 0x436715 VMOVUPD %YMM4,0x1548(%RSP,%RCX,8) |
(780) 0x43671e VMULPD %YMM6,%YMM21,%YMM4 |
(780) 0x436724 VMOVUPD %YMM4,0x1348(%RSP,%RCX,8) |
(780) 0x43672d VMULPD %YMM22,%YMM7,%YMM4 |
(780) 0x436733 VMOVUPD %YMM4,0x1148(%RSP,%RCX,8) |
(780) 0x43673c VMULPD %YMM28,%YMM8,%YMM4 |
(780) 0x436742 VMOVUPD %YMM4,0xf48(%RSP,%RCX,8) |
(780) 0x43674b VMULPD %YMM24,%YMM6,%YMM4 |
(780) 0x436751 VMOVUPD %YMM4,0xd48(%RSP,%RCX,8) |
(780) 0x43675a VMULPD %YMM22,%YMM6,%YMM4 |
(780) 0x436760 VMOVUPD %YMM4,0xb48(%RSP,%RCX,8) |
(780) 0x436769 VMULPD %YMM28,%YMM7,%YMM4 |
(780) 0x43676f VMOVUPD %YMM4,0x948(%RSP,%RCX,8) |
(780) 0x436778 VFMADD213PD (%R9,%RCX,8),%YMM28,%YMM6 |
(780) 0x43677f VMOVUPD %YMM6,(%R9,%RCX,8) |
(780) 0x436785 ADD $0x4,%RCX |
(780) 0x436789 CMP %RAX,%RCX |
(780) 0x43678c JLE 436690 |
(778) 0x436792 MOV %R10,%R14 |
(778) 0x436795 CMP %R10,%R11 |
(778) 0x436798 VMOVUPD 0x530(%RSP),%XMM0 |
(778) 0x4367a1 VMOVUPD 0x520(%RSP),%XMM1 |
(778) 0x4367aa VMOVUPD 0x920(%RSP),%YMM2 |
(778) 0x4367b3 VMOVUPD 0x510(%RSP),%XMM3 |
(778) 0x4367bc VMOVAPD %YMM9,%YMM31 |
(778) 0x4367c2 VMOVUPD 0x500(%RSP),%XMM18 |
(778) 0x4367ca VMOVUPD 0x340(%RSP),%YMM9 |
(778) 0x4367d3 VMOVUPD 0x260(%RSP),%XMM11 |
(778) 0x4367dc JNE 4367f3 |
(778) 0x4367de JMP 4369c4 |
(778) 0x4367f0 XOR %R14D,%R14D |
(778) 0x4367f3 MOV 0x120(%RSP),%RSI |
(778) 0x4367fb SUB %R14,%RSI |
(778) 0x4367fe INC %RSI |
(778) 0x436801 MOV 0x100(%RSP),%RAX |
(778) 0x436809 LEA (%RAX,%R14,8),%RDX |
(778) 0x43680d LEA 0x948(%RSP,%R14,8),%RAX |
(778) 0x436815 MOV %RAX,0x50(%RSP) |
(778) 0x43681a LEA 0xb48(%RSP,%R14,8),%RAX |
(778) 0x436822 MOV %RAX,0x48(%RSP) |
(778) 0x436827 LEA 0xd48(%RSP,%R14,8),%RAX |
(778) 0x43682f MOV %RAX,0x40(%RSP) |
(778) 0x436834 LEA 0xf48(%RSP,%R14,8),%R9 |
(778) 0x43683c LEA 0x1148(%RSP,%R14,8),%R8 |
(778) 0x436844 LEA 0x1348(%RSP,%R14,8),%R10 |
(778) 0x43684c LEA 0x1548(%RSP,%R14,8),%RBX |
(778) 0x436854 MOV 0x218(%RSP),%RCX |
(778) 0x43685c LEA (%R14,%RCX,1),%RCX |
(778) 0x436860 MOV 0x1b0(%RSP),%R12 |
(778) 0x436868 LEA (%R12,%RCX,8),%R11 |
(778) 0x43686c MOV 0x220(%RSP),%RCX |
(778) 0x436874 ADD %R14,%RCX |
(778) 0x436877 LEA (%R12,%RCX,8),%R13 |
(778) 0x43687b MOV 0x228(%RSP),%RCX |
(778) 0x436883 ADD %R14,%RCX |
(778) 0x436886 LEA (%R12,%RCX,8),%R15 |
(778) 0x43688a LEA 0x1748(%RSP,%R14,8),%RCX |
(778) 0x436892 LEA 0x1948(%RSP,%R14,8),%RDI |
(778) 0x43689a ADD 0x230(%RSP),%R14 |
(778) 0x4368a2 LEA (%R12,%R14,8),%R12 |
(778) 0x4368a6 XOR %R14D,%R14D |
(778) 0x4368a9 NOPL (%RAX) |
(776) 0x4368b0 VMOVSD (%R12,%R14,8),%XMM4 |
(776) 0x4368b6 VMOVSD (%R15,%R14,8),%XMM5 |
(776) 0x4368bc VMOVSD (%R11,%R14,8),%XMM6 |
(776) 0x4368c2 VMOVSD (%R13,%R14,8),%XMM7 |
(776) 0x4368c9 VPUNPCKLQDQ %XMM7,%XMM5,%XMM8 |
(776) 0x4368cd VMULPD %XMM10,%XMM8,%XMM8 |
(776) 0x4368d2 VSHUFPD $0x1,%XMM8,%XMM8,%XMM29 |
(776) 0x4368d9 VFMADD231SD %XMM4,%XMM0,%XMM8 |
(776) 0x4368de VFMADD231SD %XMM3,%XMM6,%XMM29 |
(776) 0x4368e4 VADDSD %XMM8,%XMM29,%XMM8 |
(776) 0x4368ea VPUNPCKLQDQ %XMM5,%XMM4,%XMM29 |
(776) 0x4368f0 VMULPD %XMM9,%XMM29,%XMM29 |
(776) 0x4368f6 VPUNPCKLQDQ %XMM6,%XMM7,%XMM30 |
(776) 0x4368fc VFMADD213PD %XMM29,%XMM18,%XMM30 |
(776) 0x436902 VSHUFPD $0x1,%XMM30,%XMM30,%XMM29 |
(776) 0x436909 VFMADD213SD %XMM30,%XMM20,%XMM29 |
(776) 0x43690f VPUNPCKLQDQ %XMM7,%XMM4,%XMM4 |
(776) 0x436913 VMULPD %XMM2,%XMM4,%XMM4 |
(776) 0x436917 VSHUFPD $0x1,%XMM4,%XMM4,%XMM7 |
(776) 0x43691c VMULSD %XMM8,%XMM27,%XMM30 |
(776) 0x436922 VMOVSD %XMM30,(%RDI,%R14,8) |
(776) 0x436929 VMULSD %XMM8,%XMM26,%XMM30 |
(776) 0x43692f VMOVSD %XMM30,(%RCX,%R14,8) |
(776) 0x436936 VMULSD %XMM25,%XMM29,%XMM30 |
(776) 0x43693c VMOVSD %XMM30,(%RBX,%R14,8) |
(776) 0x436943 VMULSD %XMM8,%XMM11,%XMM30 |
(776) 0x436949 VMOVSD %XMM30,(%R10,%R14,8) |
(776) 0x436950 VMULSD %XMM17,%XMM29,%XMM30 |
(776) 0x436956 VMOVSD %XMM30,(%R8,%R14,8) |
(776) 0x43695d VFMADD213SD %XMM4,%XMM1,%XMM5 |
(776) 0x436962 VFMADD231SD %XMM6,%XMM20,%XMM7 |
(776) 0x436968 VADDSD %XMM5,%XMM7,%XMM4 |
(776) 0x43696c VMULSD %XMM16,%XMM4,%XMM4 |
(776) 0x436972 VMOVSD %XMM4,(%R9,%R14,8) |
(776) 0x436978 VMULSD %XMM25,%XMM8,%XMM4 |
(776) 0x43697e MOV 0x40(%RSP),%RAX |
(776) 0x436983 VMOVSD %XMM4,(%RAX,%R14,8) |
(776) 0x436989 VMULSD %XMM17,%XMM8,%XMM4 |
(776) 0x43698f MOV 0x48(%RSP),%RAX |
(776) 0x436994 VMOVSD %XMM4,(%RAX,%R14,8) |
(776) 0x43699a VMULSD %XMM16,%XMM29,%XMM4 |
(776) 0x4369a0 MOV 0x50(%RSP),%RAX |
(776) 0x4369a5 VMOVSD %XMM4,(%RAX,%R14,8) |
(776) 0x4369ab VFMADD213SD (%RDX,%R14,8),%XMM16,%XMM8 |
(776) 0x4369b2 VMOVSD %XMM8,(%RDX,%R14,8) |
(776) 0x4369b8 INC %R14 |
(776) 0x4369bb CMP %R14,%RSI |
(776) 0x4369be JNE 4368b0 |
(778) 0x4369c4 MOV 0x38(%RSP),%R12 |
(778) 0x4369c9 TEST %R12,%R12 |
(778) 0x4369cc JE 436b20 |
(778) 0x4369d2 LEA -0x1(%R12),%RAX |
(778) 0x4369d7 XOR %ECX,%ECX |
(778) 0x4369d9 MOV 0x110(%RSP),%RDX |
(778) 0x4369e1 MOV 0x210(%RSP),%RSI |
(778) 0x4369e9 MOV 0x208(%RSP),%RDI |
(778) 0x4369f1 MOV 0x200(%RSP),%R8 |
(778) 0x4369f9 MOV 0x1f8(%RSP),%R9 |
(778) 0x436a01 MOV 0x1f0(%RSP),%R10 |
(778) 0x436a09 MOV 0x1e8(%RSP),%R11 |
(778) 0x436a11 MOV 0x1e0(%RSP),%RBX |
(778) 0x436a19 MOV 0x108(%RSP),%R14 |
(778) 0x436a21 NOPW %CS:(%RAX,%RAX,1) |
(779) 0x436a30 VMOVUPD 0x1948(%RSP,%RCX,8),%YMM4 |
(779) 0x436a39 VMOVUPD 0x1748(%RSP,%RCX,8),%YMM5 |
(779) 0x436a42 VMOVUPD 0x1548(%RSP,%RCX,8),%YMM6 |
(779) 0x436a4b VMOVUPD 0x1348(%RSP,%RCX,8),%YMM7 |
(779) 0x436a54 VMOVUPD 0x1148(%RSP,%RCX,8),%YMM8 |
(779) 0x436a5d VMOVUPD 0xf48(%RSP,%RCX,8),%YMM29 |
(779) 0x436a68 VMOVUPD 0xd48(%RSP,%RCX,8),%YMM30 |
(779) 0x436a73 VMOVUPD 0xb48(%RSP,%RCX,8),%YMM9 |
(779) 0x436a7c VMOVUPD 0x948(%RSP,%RCX,8),%YMM10 |
(779) 0x436a85 VADDPD (%RDX,%RCX,8),%YMM4,%YMM4 |
(779) 0x436a8a VMOVUPD %YMM4,(%RDX,%RCX,8) |
(779) 0x436a8f VADDPD (%RSI,%RCX,8),%YMM5,%YMM4 |
(779) 0x436a94 VMOVUPD %YMM4,(%RSI,%RCX,8) |
(779) 0x436a99 VADDPD (%R8,%RCX,8),%YMM6,%YMM4 |
(779) 0x436a9f VMOVUPD %YMM4,(%R8,%RCX,8) |
(779) 0x436aa5 VADDPD (%R10,%RCX,8),%YMM7,%YMM4 |
(779) 0x436aab VMOVUPD %YMM4,(%R10,%RCX,8) |
(779) 0x436ab1 VADDPD (%R11,%RCX,8),%YMM8,%YMM4 |
(779) 0x436ab7 VMOVUPD %YMM4,(%R11,%RCX,8) |
(779) 0x436abd VADDPD (%RBX,%RCX,8),%YMM29,%YMM4 |
(779) 0x436ac4 VMOVUPD %YMM4,(%RBX,%RCX,8) |
(779) 0x436ac9 VADDPD (%R14,%RCX,8),%YMM30,%YMM4 |
(779) 0x436ad0 VMOVUPD %YMM4,(%R14,%RCX,8) |
(779) 0x436ad6 VADDPD (%RDI,%RCX,8),%YMM9,%YMM4 |
(779) 0x436adb VMOVUPD %YMM4,(%RDI,%RCX,8) |
(779) 0x436ae0 VADDPD (%R9,%RCX,8),%YMM10,%YMM4 |
(779) 0x436ae6 VMOVUPD %YMM4,(%R9,%RCX,8) |
(779) 0x436aec ADD $0x4,%RCX |
(779) 0x436af0 CMP %RAX,%RCX |
(779) 0x436af3 JLE 436a30 |
(778) 0x436af9 CMP %R12,0x118(%RSP) |
(778) 0x436b01 MOV 0x120(%RSP),%RAX |
(778) 0x436b09 JE 436460 |
(778) 0x436b0f JMP 436b2b |
(778) 0x436b20 XOR %R12D,%R12D |
(778) 0x436b23 MOV 0x120(%RSP),%RAX |
(778) 0x436b2b SUB %R12,%RAX |
(778) 0x436b2e INC %RAX |
(778) 0x436b31 MOV %RAX,0x120(%RSP) |
(778) 0x436b39 MOV 0x108(%RSP),%RAX |
(778) 0x436b41 LEA (%RAX,%R12,8),%RAX |
(778) 0x436b45 MOV 0x238(%RSP),%RCX |
(778) 0x436b4d ADD %R12,%RCX |
(778) 0x436b50 MOV 0x18(%RSP),%R11 |
(778) 0x436b55 LEA (%R11,%RCX,8),%RCX |
(778) 0x436b59 MOV 0x240(%RSP),%RDX |
(778) 0x436b61 LEA (%R12,%RDX,1),%RDX |
(778) 0x436b65 LEA (%R11,%RDX,8),%RDX |
(778) 0x436b69 MOV 0x248(%RSP),%RSI |
(778) 0x436b71 LEA (%R12,%RSI,1),%RSI |
(778) 0x436b75 LEA (%R11,%RSI,8),%RSI |
(778) 0x436b79 MOV 0x250(%RSP),%RDI |
(778) 0x436b81 LEA (%R12,%RDI,1),%R8 |
(778) 0x436b85 MOV 0x28(%RSP),%R9 |
(778) 0x436b8a LEA (%R9,%R8,8),%RDI |
(778) 0x436b8e LEA (%R11,%R8,8),%R8 |
(778) 0x436b92 MOV 0x258(%RSP),%R10 |
(778) 0x436b9a LEA (%R12,%R10,1),%R10 |
(778) 0x436b9e LEA (%R9,%R10,8),%R9 |
(778) 0x436ba2 LEA (%R11,%R10,8),%R10 |
(778) 0x436ba6 MOV 0x110(%RSP),%R11 |
(778) 0x436bae LEA (%R11,%R12,8),%R11 |
(778) 0x436bb2 LEA 0x948(%RSP,%R12,8),%RBX |
(778) 0x436bba MOV %RBX,0x50(%RSP) |
(778) 0x436bbf LEA 0xb48(%RSP,%R12,8),%RBX |
(778) 0x436bc7 MOV %RBX,0x48(%RSP) |
(778) 0x436bcc LEA 0xd48(%RSP,%R12,8),%RBX |
(778) 0x436bd4 MOV %RBX,0x40(%RSP) |
(778) 0x436bd9 LEA 0xf48(%RSP,%R12,8),%RBX |
(778) 0x436be1 MOV %RBX,0x38(%RSP) |
(778) 0x436be6 LEA 0x1148(%RSP,%R12,8),%RBX |
(778) 0x436bee MOV %RBX,0x118(%RSP) |
(778) 0x436bf6 LEA 0x1348(%RSP,%R12,8),%RBX |
(778) 0x436bfe MOV %RBX,0x338(%RSP) |
(778) 0x436c06 LEA 0x1548(%RSP,%R12,8),%R14 |
(778) 0x436c0e LEA 0x1748(%RSP,%R12,8),%R15 |
(778) 0x436c16 LEA 0x1948(%RSP,%R12,8),%R12 |
(778) 0x436c1e XOR %R13D,%R13D |
(778) 0x436c21 NOPW %CS:(%RAX,%RAX,1) |
(777) 0x436c30 VMOVSD (%R12,%R13,8),%XMM4 |
(777) 0x436c36 VMOVSD (%R15,%R13,8),%XMM5 |
(777) 0x436c3c VMOVSD (%R14,%R13,8),%XMM6 |
(777) 0x436c42 MOV 0x338(%RSP),%RBX |
(777) 0x436c4a VMOVSD (%RBX,%R13,8),%XMM7 |
(777) 0x436c50 MOV 0x118(%RSP),%RBX |
(777) 0x436c58 VMOVSD (%RBX,%R13,8),%XMM8 |
(777) 0x436c5e MOV 0x38(%RSP),%RBX |
(777) 0x436c63 VMOVSD (%RBX,%R13,8),%XMM9 |
(777) 0x436c69 MOV 0x40(%RSP),%RBX |
(777) 0x436c6e VMOVSD (%RBX,%R13,8),%XMM10 |
(777) 0x436c74 MOV 0x48(%RSP),%RBX |
(777) 0x436c79 VMOVSD (%RBX,%R13,8),%XMM29 |
(777) 0x436c80 MOV 0x50(%RSP),%RBX |
(777) 0x436c85 VMOVSD (%RBX,%R13,8),%XMM30 |
(777) 0x436c8c VADDSD (%R11,%R13,8),%XMM4,%XMM4 |
(777) 0x436c92 VMOVSD %XMM4,(%R11,%R13,8) |
(777) 0x436c98 VADDSD (%R10,%R13,8),%XMM5,%XMM4 |
(777) 0x436c9e VMOVSD %XMM4,(%R10,%R13,8) |
(777) 0x436ca4 VADDSD (%R8,%R13,8),%XMM6,%XMM4 |
(777) 0x436caa VMOVSD %XMM4,(%R8,%R13,8) |
(777) 0x436cb0 VADDSD (%RSI,%R13,8),%XMM7,%XMM4 |
(777) 0x436cb6 VMOVSD %XMM4,(%RSI,%R13,8) |
(777) 0x436cbc VADDSD (%RDX,%R13,8),%XMM8,%XMM4 |
(777) 0x436cc2 VMOVSD %XMM4,(%RDX,%R13,8) |
(777) 0x436cc8 VADDSD (%RCX,%R13,8),%XMM9,%XMM4 |
(777) 0x436cce VMOVSD %XMM4,(%RCX,%R13,8) |
(777) 0x436cd4 VADDSD (%RAX,%R13,8),%XMM10,%XMM4 |
(777) 0x436cda VMOVSD %XMM4,(%RAX,%R13,8) |
(777) 0x436ce0 VADDSD (%R9,%R13,8),%XMM29,%XMM4 |
(777) 0x436ce7 VMOVSD %XMM4,(%R9,%R13,8) |
(777) 0x436ced VADDSD (%RDI,%R13,8),%XMM30,%XMM4 |
(777) 0x436cf4 VMOVSD %XMM4,(%RDI,%R13,8) |
(777) 0x436cfa INC %R13 |
(777) 0x436cfd CMP %R13,0x120(%RSP) |
(777) 0x436d05 JNE 436c30 |
(778) 0x436d0b JMP 436460 |
0x436d10 TEST %R10D,%R10D |
0x436d13 JLE 4361e8 |
0x436d19 TEST %RCX,%RCX |
0x436d1c JE 436f62 |
0x436d22 VBROADCASTSD %XMM27,%YMM4 |
0x436d28 VMOVUPD %YMM4,0x120(%RSP) |
0x436d31 VBROADCASTSD %XMM26,%YMM21 |
0x436d37 VBROADCASTSD %XMM25,%YMM22 |
0x436d3d VBROADCASTSD 0x260(%RSP),%YMM23 |
0x436d45 VBROADCASTSD %XMM17,%YMM24 |
0x436d4b VBROADCASTSD %XMM16,%YMM28 |
0x436d51 XOR %EAX,%EAX |
0x436d53 MOV 0xd8(%RSP),%RDX |
0x436d5b MOV 0xc0(%RSP),%R15 |
0x436d63 MOV 0xd0(%RSP),%R13 |
0x436d6b MOV 0xf8(%RSP),%R14 |
0x436d73 MOV 0xf0(%RSP),%R11 |
0x436d7b MOV 0xe8(%RSP),%R10 |
0x436d83 MOV 0xe0(%RSP),%RDI |
0x436d8b VMOVAPD %YMM31,%YMM13 |
0x436d91 NOPW %CS:(%RAX,%RAX,1) |
(775) 0x436da0 VMOVUPD (%R14,%RAX,8),%YMM4 |
(775) 0x436da6 VMOVUPD (%R11,%RAX,8),%YMM5 |
(775) 0x436dac VMOVUPD (%R10,%RAX,8),%YMM6 |
(775) 0x436db2 VMOVUPD (%RDI,%RAX,8),%YMM7 |
(775) 0x436db7 VMULPD %YMM8,%YMM4,%YMM19 |
(775) 0x436dbd VFMADD231PD %YMM14,%YMM5,%YMM19 |
(775) 0x436dc3 VFMADD231PD %YMM15,%YMM6,%YMM19 |
(775) 0x436dc9 VFMADD231PD %YMM7,%YMM30,%YMM19 |
(775) 0x436dcf VMOVAPD %YMM30,%YMM31 |
(775) 0x436dd5 VMOVAPD %YMM15,%YMM30 |
(775) 0x436ddb VMOVAPD %YMM14,%YMM15 |
(775) 0x436de0 VMOVAPD %YMM8,%YMM14 |
(775) 0x436de5 VMULPD %YMM29,%YMM4,%YMM8 |
(775) 0x436deb VMULPD %YMM5,%YMM11,%YMM9 |
(775) 0x436def VFMADD231PD 0x360(%RSP),%YMM7,%YMM9 |
(775) 0x436df9 VFMADD231PD 0x380(%RSP),%YMM6,%YMM8 |
(775) 0x436e03 VFMADD231PD %YMM9,%YMM12,%YMM8 |
(775) 0x436e08 LEA (%R8,%RAX,8),%RCX |
(775) 0x436e0c VMOVUPD (%R8,%RAX,8),%YMM9 |
(775) 0x436e12 VFMADD231PD 0x120(%RSP),%YMM19,%YMM9 |
(775) 0x436e1a VMOVUPD %YMM9,(%R8,%RAX,8) |
(775) 0x436e20 VMOVUPD (%R12,%RCX,1),%YMM9 |
(775) 0x436e26 VFMADD231PD %YMM19,%YMM21,%YMM9 |
(775) 0x436e2c VMOVUPD %YMM9,(%R12,%RCX,1) |
(775) 0x436e32 VMULPD 0x460(%RSP),%YMM4,%YMM4 |
(775) 0x436e3b LEA (%RCX,%R12,1),%RCX |
(775) 0x436e3f VMOVUPD (%R12,%RCX,1),%YMM9 |
(775) 0x436e45 VFMADD231PD %YMM22,%YMM8,%YMM9 |
(775) 0x436e4b VMOVUPD %YMM9,(%R12,%RCX,1) |
(775) 0x436e51 VFMADD231PD %YMM5,%YMM13,%YMM4 |
(775) 0x436e56 LEA (%RCX,%R12,1),%RCX |
(775) 0x436e5a VMOVUPD (%R12,%RCX,1),%YMM5 |
(775) 0x436e60 VFMADD231PD %YMM19,%YMM23,%YMM5 |
(775) 0x436e66 VMOVUPD %YMM5,(%R12,%RCX,1) |
(775) 0x436e6c VFMADD231PD 0x440(%RSP),%YMM6,%YMM4 |
(775) 0x436e76 LEA (%RCX,%R12,1),%RCX |
(775) 0x436e7a VMOVUPD (%R12,%RCX,1),%YMM5 |
(775) 0x436e80 VFMADD231PD %YMM24,%YMM8,%YMM5 |
(775) 0x436e86 VMOVUPD %YMM5,(%R12,%RCX,1) |
(775) 0x436e8c VFMADD231PD %YMM7,%YMM12,%YMM4 |
(775) 0x436e91 LEA (%RCX,%R12,1),%RCX |
(775) 0x436e95 VFMADD213PD (%R12,%RCX,1),%YMM28,%YMM4 |
(775) 0x436e9c VMOVUPD %YMM4,(%R12,%RCX,1) |
(775) 0x436ea2 MOV 0x30(%RSP),%RCX |
(775) 0x436ea7 VMOVUPD (%R9,%RAX,8),%YMM4 |
(775) 0x436ead VFMADD231PD %YMM22,%YMM19,%YMM4 |
(775) 0x436eb3 VMOVUPD %YMM4,(%R9,%RAX,8) |
(775) 0x436eb9 VMOVUPD (%R13,%RAX,8),%YMM4 |
(775) 0x436ec0 VFMADD231PD %YMM24,%YMM19,%YMM4 |
(775) 0x436ec6 VMOVUPD %YMM4,(%R13,%RAX,8) |
(775) 0x436ecd VFMADD213PD (%R15,%RAX,8),%YMM28,%YMM8 |
(775) 0x436ed4 VMOVUPD %YMM8,(%R15,%RAX,8) |
(775) 0x436eda VMOVAPD %YMM14,%YMM8 |
(775) 0x436edf VMOVAPD %YMM15,%YMM14 |
(775) 0x436ee4 VMOVAPD %YMM30,%YMM15 |
(775) 0x436eea VMOVAPD %YMM31,%YMM30 |
(775) 0x436ef0 VFMADD213PD (%RDX,%RAX,8),%YMM28,%YMM19 |
(775) 0x436ef7 VMOVUPD %YMM19,(%RDX,%RAX,8) |
(775) 0x436efe ADD $0x4,%RAX |
(775) 0x436f02 CMP %RCX,%RAX |
(775) 0x436f05 JL 436da0 |
0x436f0b MOV %RCX,%RAX |
0x436f0e MOV 0x90(%RSP),%R10 |
0x436f16 CMP %R10,%RCX |
0x436f19 VMOVUPD 0x340(%RSP),%YMM9 |
0x436f22 VMOVUPD 0x380(%RSP),%YMM7 |
0x436f2b MOV 0x1b8(%RSP),%R13 |
0x436f33 VMOVUPD 0x360(%RSP),%YMM5 |
0x436f3c VMOVSD 0x180(%RSP),%XMM28 |
0x436f44 VMOVSD 0x178(%RSP),%XMM24 |
0x436f4c MOV 0x88(%RSP),%R11 |
0x436f54 VMOVAPD %YMM13,%YMM31 |
0x436f5a JE 4361e8 |
0x436f60 JMP 436f64 |
0x436f62 XOR %EAX,%EAX |
0x436f64 MOV %R10,%RCX |
0x436f67 SUB %RAX,%RCX |
0x436f6a VPBROADCASTQ %RCX,%YMM4 |
0x436f70 VPCMPNLEUQ 0xad945(%RIP),%YMM4,%K1 |
0x436f7b KORTESTB %K1,%K1 |
0x436f7f JE 43731a |
0x436f85 MOV 0x150(%RSP),%RCX |
0x436f8d ADD %RAX,%RCX |
0x436f90 MOV 0xc8(%RSP),%RDX |
0x436f98 IMUL 0x1a8(%RSP),%RDX |
0x436fa1 ADD 0x2f8(%RSP),%RDX |
0x436fa9 ADD %RDX,%RCX |
0x436fac MOV 0x1b0(%RSP),%RSI |
0x436fb4 VMOVUPD (%RSI,%RCX,8),%YMM4{%K1}{z} |
0x436fbb MOV 0x2e8(%RSP),%RCX |
0x436fc3 ADD %RAX,%RCX |
0x436fc6 ADD %RDX,%RCX |
0x436fc9 VMOVAPD %YMM5,%YMM13 |
0x436fcd VMOVUPD (%RSI,%RCX,8),%YMM5{%K1}{z} |
0x436fd4 VMOVUPD 0x760(%RSP),%YMM23 |
0x436fdc VMOVAPD %YMM4,%YMM23{%K1} |
0x436fe2 VMOVUPD 0x780(%RSP),%YMM22 |
0x436fea VMOVAPD %YMM5,%YMM22{%K1} |
0x436ff0 MOV 0x2e0(%RSP),%RCX |
0x436ff8 ADD %RAX,%RCX |
0x436ffb ADD %RDX,%RCX |
0x436ffe VMOVUPD (%RSI,%RCX,8),%YMM4{%K1}{z} |
0x437005 VMOVUPD 0x7a0(%RSP),%YMM10 |
0x43700e VMOVAPD %YMM4,%YMM10{%K1} |
0x437014 MOV 0x2f0(%RSP),%RCX |
0x43701c ADD %RAX,%RCX |
0x43701f ADD %RCX,%RDX |
0x437022 VMOVUPD (%RSI,%RDX,8),%YMM4{%K1}{z} |
0x437029 VMOVUPD 0x7c0(%RSP),%YMM9 |
0x437032 VMOVAPD %YMM4,%YMM9{%K1} |
0x437038 VMULPD %YMM8,%YMM23,%YMM19 |
0x43703e VFMADD231PD %YMM14,%YMM22,%YMM19 |
0x437044 VFMADD231PD %YMM15,%YMM10,%YMM19 |
0x43704a VFMADD231PD %YMM9,%YMM30,%YMM19 |
0x437050 VMOVUPD (%R8,%RAX,8),%YMM4{%K1}{z} |
0x437057 VMULPD %YMM29,%YMM23,%YMM21 |
0x43705d VMOVUPD 0x7e0(%RSP),%YMM6 |
0x437066 VMOVAPD %YMM4,%YMM6{%K1} |
0x43706c VMULPD %YMM11,%YMM22,%YMM4 |
0x437072 VFMADD231PD %YMM13,%YMM9,%YMM4 |
0x437077 VFMADD231PD %YMM7,%YMM10,%YMM21 |
0x43707d VBROADCASTSD %XMM27,%YMM5 |
0x437083 VMOVUPD %YMM6,0x7e0(%RSP) |
0x43708c VFMADD213PD %YMM6,%YMM19,%YMM5 |
0x437092 VMOVUPD %YMM5,(%R8,%RAX,8){%K1} |
0x437099 MOV 0x80(%RSP),%RCX |
0x4370a1 LEA (%RCX,%RAX,1),%RCX |
0x4370a5 VMOVUPD (%R8,%RCX,8),%YMM5{%K1}{z} |
0x4370ac VFMADD231PD %YMM4,%YMM12,%YMM21 |
0x4370b2 VMOVUPD 0x800(%RSP),%YMM6 |
0x4370bb VMOVAPD %YMM5,%YMM6{%K1} |
0x4370c1 VBROADCASTSD %XMM26,%YMM4 |
0x4370c7 VMOVUPD %YMM6,0x800(%RSP) |
0x4370d0 VFMADD213PD %YMM6,%YMM19,%YMM4 |
0x4370d6 VMOVUPD %YMM4,(%R8,%RCX,8){%K1} |
0x4370dd VBROADCASTSD %XMM25,%YMM4 |
0x4370e3 MOV 0x78(%RSP),%RDX |
0x4370e8 LEA (%RDX,%RAX,1),%RDX |
0x4370ec VMOVUPD (%R8,%RDX,8),%YMM5{%K1}{z} |
0x4370f3 VMOVUPD 0x820(%RSP),%YMM6 |
0x4370fc VMOVAPD %YMM5,%YMM6{%K1} |
0x437102 VMOVAPD %YMM4,%YMM5 |
0x437106 VMOVUPD %YMM6,0x820(%RSP) |
0x43710f VFMADD213PD %YMM6,%YMM21,%YMM5 |
0x437115 VMOVUPD %YMM5,(%R8,%RDX,8){%K1} |
0x43711c MOV 0x68(%RSP),%RSI |
0x437121 LEA (%RSI,%RAX,1),%RSI |
0x437125 VMOVUPD (%R8,%RSI,8),%YMM5{%K1}{z} |
0x43712c VMOVUPD 0x840(%RSP),%YMM6 |
0x437135 VMOVAPD %YMM5,%YMM6{%K1} |
0x43713b VBROADCASTSD 0x260(%RSP),%YMM5 |
0x437145 VMOVUPD %YMM6,0x840(%RSP) |
0x43714e VFMADD213PD %YMM6,%YMM19,%YMM5 |
0x437154 VMOVUPD %YMM5,(%R8,%RSI,8){%K1} |
0x43715b VBROADCASTSD %XMM17,%YMM5 |
0x437161 MOV 0x70(%RSP),%RSI |
0x437166 LEA (%RSI,%RAX,1),%RSI |
0x43716a VMOVUPD (%R8,%RSI,8),%YMM6{%K1}{z} |
0x437171 VMOVAPD %YMM11,%YMM25 |
0x437177 VMOVAPD %YMM7,%YMM11 |
0x43717b VMOVUPD 0x860(%RSP),%YMM7 |
0x437184 VMOVAPD %YMM6,%YMM7{%K1} |
0x43718a VMOVAPD %YMM5,%YMM6 |
0x43718e VMOVUPD %YMM7,0x860(%RSP) |
0x437197 VFMADD213PD %YMM7,%YMM21,%YMM6 |
0x43719d VMOVUPD %YMM6,(%R8,%RSI,8){%K1} |
0x4371a4 MOV 0x60(%RSP),%RSI |
0x4371a9 LEA (%RSI,%RAX,1),%RSI |
0x4371ad VMOVUPD (%R8,%RSI,8),%YMM6{%K1}{z} |
0x4371b4 VMOVAPD %YMM15,%YMM17 |
0x4371ba VMOVAPD %YMM14,%YMM15 |
0x4371bf VMOVAPD %YMM8,%YMM14 |
0x4371c4 VMOVUPD 0x880(%RSP),%YMM8 |
0x4371cd VMOVAPD %YMM6,%YMM8{%K1} |
0x4371d3 VMOVUPD %YMM23,0x760(%RSP) |
0x4371db VMULPD 0x460(%RSP),%YMM23,%YMM6 |
0x4371e3 VMOVUPD %YMM22,0x780(%RSP) |
0x4371eb VFMADD231PD %YMM31,%YMM22,%YMM6 |
0x4371f1 VMOVUPD %YMM10,0x7a0(%RSP) |
0x4371fa VFMADD231PD 0x440(%RSP),%YMM10,%YMM6 |
0x437204 VMOVUPD 0x420(%RSP),%YMM10 |
0x43720d VMOVUPD %YMM9,0x7c0(%RSP) |
0x437216 VFMADD231PD %YMM12,%YMM9,%YMM6 |
0x43721b VMOVUPD 0x340(%RSP),%YMM9 |
0x437224 VBROADCASTSD %XMM16,%YMM7 |
0x43722a VMOVUPD %YMM8,0x880(%RSP) |
0x437233 VFMADD213PD %YMM8,%YMM7,%YMM6 |
0x437238 VMOVUPD %YMM6,(%R8,%RSI,8){%K1} |
0x43723f MOV 0x20(%RSP),%RSI |
0x437244 VMOVUPD (%R9,%RAX,8),%YMM6{%K1}{z} |
0x43724b VMOVUPD 0x8a0(%RSP),%YMM8 |
0x437254 VMOVAPD %YMM6,%YMM8{%K1} |
0x43725a VMOVUPD %YMM8,0x8a0(%RSP) |
0x437263 VFMADD213PD %YMM8,%YMM19,%YMM4 |
0x437269 VMOVAPD %YMM14,%YMM8 |
0x43726e VMOVAPD %YMM15,%YMM14 |
0x437273 VMOVAPD %YMM17,%YMM15 |
0x437279 VMOVUPD %YMM4,(%R9,%RAX,8){%K1} |
0x437280 VMOVUPD (%R9,%RCX,8),%YMM4{%K1}{z} |
0x437287 VMOVUPD 0x8c0(%RSP),%YMM6 |
0x437290 VMOVAPD %YMM4,%YMM6{%K1} |
0x437296 VMOVUPD %YMM6,0x8c0(%RSP) |
0x43729f VFMADD213PD %YMM6,%YMM19,%YMM5 |
0x4372a5 VMOVUPD %YMM5,(%R9,%RCX,8){%K1} |
0x4372ac VMOVUPD (%R9,%RDX,8),%YMM4{%K1}{z} |
0x4372b3 VMOVUPD 0x8e0(%RSP),%YMM5 |
0x4372bc VMOVAPD %YMM4,%YMM5{%K1} |
0x4372c2 VMOVUPD %YMM5,0x8e0(%RSP) |
0x4372cb VFMADD213PD %YMM5,%YMM7,%YMM21 |
0x4372d1 VMOVUPD %YMM21,(%R9,%RDX,8){%K1} |
0x4372d8 MOV 0xd8(%RSP),%RCX |
0x4372e0 VMOVUPD (%RCX,%RAX,8),%YMM4{%K1}{z} |
0x4372e7 VMOVUPD 0x900(%RSP),%YMM5 |
0x4372f0 VMOVAPD %YMM4,%YMM5{%K1} |
0x4372f6 VMOVUPD %YMM5,0x900(%RSP) |
0x4372ff VFMADD213PD %YMM5,%YMM7,%YMM19 |
0x437305 VMOVAPD %YMM13,%YMM5 |
0x437309 VMOVAPD %YMM11,%YMM7 |
0x43730d VMOVAPD %YMM25,%YMM11 |
0x437313 VMOVUPD %YMM19,(%RCX,%RAX,8){%K1} |
0x43731a MOV 0x30(%RSP),%RCX |
0x43731f JMP 4361e8 |
/scratch_na/users/xoserete/qaas_runs/171-284-5202/intel/miniqmc/build/miniqmc/src/Numerics/Spline2/MultiBsplineRef.hpp: 227 - 262 |
-------------------------------------------------------------------------------- |
227: for (int j = 0; j < 4; j++) |
[...] |
234: const T pre20 = d2a[i] * b[j]; |
235: const T pre10 = da[i] * b[j]; |
236: const T pre00 = a[i] * b[j]; |
237: const T pre11 = da[i] * db[j]; |
238: const T pre01 = a[i] * db[j]; |
239: const T pre02 = a[i] * d2b[j]; |
240: |
241: const int iSplitPoint = num_splines; |
242: for (int n = 0; n < iSplitPoint; n++) |
243: { |
244: T coefsv = coefs[n]; |
245: T coefsvzs = coefszs[n]; |
246: T coefsv2zs = coefs2zs[n]; |
247: T coefsv3zs = coefs3zs[n]; |
248: |
249: T sum0 = c[0] * coefsv + c[1] * coefsvzs + c[2] * coefsv2zs + c[3] * coefsv3zs; |
250: T sum1 = dc[0] * coefsv + dc[1] * coefsvzs + dc[2] * coefsv2zs + dc[3] * coefsv3zs; |
251: T sum2 = d2c[0] * coefsv + d2c[1] * coefsvzs + d2c[2] * coefsv2zs + d2c[3] * coefsv3zs; |
252: |
253: hxx[n] += pre20 * sum0; |
254: hxy[n] += pre11 * sum0; |
255: hxz[n] += pre10 * sum1; |
256: hyy[n] += pre02 * sum0; |
257: hyz[n] += pre01 * sum1; |
258: hzz[n] += pre00 * sum2; |
259: gx[n] += pre10 * sum0; |
260: gy[n] += pre01 * sum0; |
261: gz[n] += pre00 * sum1; |
262: vals[n] += pre00 * sum0; |
/scratch_na/users/xoserete/qaas_runs/171-284-5202/intel/miniqmc/build/miniqmc/src/Numerics/OhmmsPETE/TinyVector.h: 61 - 61 |
-------------------------------------------------------------------------------- |
61: for (size_t d = 0; d < D; ++d) |
Path / |
Metric | Value |
---|---|
CQA speedup if no scalar integer | 1.37 |
CQA speedup if FP arith vectorized | 1.63 |
CQA speedup if fully vectorized | 4.42 |
CQA speedup if no inter-iteration dependency | NA |
CQA speedup if next bottleneck killed | 1.22 |
Bottlenecks | micro-operation queue, |
Function | miniqmcreference::einspline_spo_ref |
Source | MultiBsplineRef.hpp:227-227,MultiBsplineRef.hpp:234-239,MultiBsplineRef.hpp:242-262,TinyVector.h:61-61 |
Source loop unroll info | NA |
Source loop unroll confidence level | NA |
Unroll/vectorization loop type | NA |
Unroll factor | NA |
CQA cycles | 48.00 |
CQA cycles if no scalar integer | 35.00 |
CQA cycles if FP arith vectorized | 29.39 |
CQA cycles if fully vectorized | 10.85 |
Front-end cycles | 48.00 |
DIV/SQRT cycles | 16.50 |
P0 cycles | 16.00 |
P1 cycles | 39.33 |
P2 cycles | 39.33 |
P3 cycles | 28.50 |
P4 cycles | 17.00 |
P5 cycles | 14.70 |
P6 cycles | 28.50 |
P7 cycles | 28.50 |
P8 cycles | 28.50 |
P9 cycles | 14.80 |
P10 cycles | 39.33 |
P11 cycles | 0.00 |
Inter-iter dependencies cycles | NA |
FE+BE cycles (UFS) | 49.65 |
Stall cycles (UFS) | 0.00 |
Nb insns | 280.00 |
Nb uops | 287.00 |
Nb loads | 119.00 |
Nb stores | 57.00 |
Nb stack references | 89.00 |
FLOP/cycle | 3.63 |
Nb FLOP add-sub | 0.00 |
Nb FLOP mul | 22.00 |
Nb FLOP fma | 76.00 |
Nb FLOP div | 0.00 |
Nb FLOP rcp | 0.00 |
Nb FLOP sqrt | 0.00 |
Nb FLOP rsqrt | 0.00 |
Bytes/cycle | 65.67 |
Bytes prefetched | 0.00 |
Bytes loaded | 2088.00 |
Bytes stored | 1064.00 |
Stride 0 | NA |
Stride 1 | NA |
Stride n | NA |
Stride unknown | NA |
Stride indirect | NA |
Vectorization ratio all | 60.98 |
Vectorization ratio load | 63.16 |
Vectorization ratio store | 45.61 |
Vectorization ratio mul | 40.00 |
Vectorization ratio add_sub | NA |
Vectorization ratio fma | 100.00 |
Vectorization ratio div_sqrt | NA |
Vectorization ratio other | 54.39 |
Vector-efficiency ratio all | 35.00 |
Vector-efficiency ratio load | 35.86 |
Vector-efficiency ratio store | 29.17 |
Vector-efficiency ratio mul | 27.50 |
Vector-efficiency ratio add_sub | NA |
Vector-efficiency ratio fma | 50.00 |
Vector-efficiency ratio div_sqrt | NA |
Vector-efficiency ratio other | 32.46 |
Metric | Value |
---|---|
CQA speedup if no scalar integer | 1.37 |
CQA speedup if FP arith vectorized | 1.63 |
CQA speedup if fully vectorized | 4.42 |
CQA speedup if no inter-iteration dependency | NA |
CQA speedup if next bottleneck killed | 1.22 |
Bottlenecks | micro-operation queue, |
Function | miniqmcreference::einspline_spo_ref |
Source | MultiBsplineRef.hpp:227-227,MultiBsplineRef.hpp:234-239,MultiBsplineRef.hpp:242-262,TinyVector.h:61-61 |
Source loop unroll info | NA |
Source loop unroll confidence level | NA |
Unroll/vectorization loop type | NA |
Unroll factor | NA |
CQA cycles | 48.00 |
CQA cycles if no scalar integer | 35.00 |
CQA cycles if FP arith vectorized | 29.39 |
CQA cycles if fully vectorized | 10.85 |
Front-end cycles | 48.00 |
DIV/SQRT cycles | 16.50 |
P0 cycles | 16.00 |
P1 cycles | 39.33 |
P2 cycles | 39.33 |
P3 cycles | 28.50 |
P4 cycles | 17.00 |
P5 cycles | 14.70 |
P6 cycles | 28.50 |
P7 cycles | 28.50 |
P8 cycles | 28.50 |
P9 cycles | 14.80 |
P10 cycles | 39.33 |
P11 cycles | 0.00 |
Inter-iter dependencies cycles | NA |
FE+BE cycles (UFS) | 49.65 |
Stall cycles (UFS) | 0.00 |
Nb insns | 280.00 |
Nb uops | 287.00 |
Nb loads | 119.00 |
Nb stores | 57.00 |
Nb stack references | 89.00 |
FLOP/cycle | 3.63 |
Nb FLOP add-sub | 0.00 |
Nb FLOP mul | 22.00 |
Nb FLOP fma | 76.00 |
Nb FLOP div | 0.00 |
Nb FLOP rcp | 0.00 |
Nb FLOP sqrt | 0.00 |
Nb FLOP rsqrt | 0.00 |
Bytes/cycle | 65.67 |
Bytes prefetched | 0.00 |
Bytes loaded | 2088.00 |
Bytes stored | 1064.00 |
Stride 0 | NA |
Stride 1 | NA |
Stride n | NA |
Stride unknown | NA |
Stride indirect | NA |
Vectorization ratio all | 60.98 |
Vectorization ratio load | 63.16 |
Vectorization ratio store | 45.61 |
Vectorization ratio mul | 40.00 |
Vectorization ratio add_sub | NA |
Vectorization ratio fma | 100.00 |
Vectorization ratio div_sqrt | NA |
Vectorization ratio other | 54.39 |
Vector-efficiency ratio all | 35.00 |
Vector-efficiency ratio load | 35.86 |
Vector-efficiency ratio store | 29.17 |
Vector-efficiency ratio mul | 27.50 |
Vector-efficiency ratio add_sub | NA |
Vector-efficiency ratio fma | 50.00 |
Vector-efficiency ratio div_sqrt | NA |
Vector-efficiency ratio other | 32.46 |
Path / |
Function | miniqmcreference::einspline_spo_ref |
Source file and lines | MultiBsplineRef.hpp:227-262 |
Module | exec |
nb instructions | 280 |
nb uops | 287 |
loop length | 1919 |
used x86 registers | 14 |
used mmx registers | 0 |
used xmm registers | 9 |
used ymm registers | 23 |
used zmm registers | 0 |
nb stack references | 89 |
micro-operation queue | 48.00 cycles |
front end | 48.00 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 16.50 | 16.00 | 39.33 | 39.33 | 28.50 | 17.00 | 14.70 | 28.50 | 28.50 | 28.50 | 14.80 | 39.33 |
cycles | 16.50 | 16.00 | 39.33 | 39.33 | 28.50 | 17.00 | 14.70 | 28.50 | 28.50 | 28.50 | 14.80 | 39.33 |
Cycles executing div or sqrt instructions | NA |
FE+BE cycles | 49.65 |
Stall cycles | 0.00 |
Front-end | 48.00 |
Dispatch | 39.33 |
Overall L1 | 48.00 |
all | 2% |
load | 5% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 11% |
all | 80% |
load | 82% |
store | 100% |
mul | 40% |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | 100% |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 62% |
all | 60% |
load | 63% |
store | 45% |
mul | 40% |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | 100% |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 54% |
all | 12% |
load | 14% |
store | 12% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 13% |
all | 42% |
load | 42% |
store | 49% |
mul | 27% |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | 50% |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 35% |
all | 35% |
load | 35% |
store | 29% |
mul | 27% |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | 50% |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 32% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
MOV 0x30(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD 0x620(%RSP),%YMM8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x600(%RSP),%YMM14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5e0(%RSP),%YMM15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5c0(%RSP),%YMM30 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5a0(%RSP),%YMM29 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x580(%RSP),%YMM11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x380(%RSP),%YMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x360(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVSD 0x180(%RSP),%XMM28 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD 0x178(%RSP),%XMM24 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x90(%RSP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x18(%RSP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x1a8(%RSP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA 0x1(%RDI),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x308(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RDX,0xe0(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0xe8(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0xf0(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0xf8(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
MOV 0xc8(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RDX,0x188(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0x190(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0x198(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0x1a0(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
CMP $0x3,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JE 436030 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xb90> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VMOVSD 0x540(%RSP,%RAX,8),%XMM16 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMULSD %XMM28,%XMM16,%XMM27 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMULSD %XMM24,%XMM16,%XMM25 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVSD 0x310(%RSP),%XMM4 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMULSD %XMM4,%XMM16,%XMM16 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVSD 0x3e0(%RSP,%RAX,8),%XMM17 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMULSD %XMM24,%XMM17,%XMM26 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMULSD %XMM4,%XMM17,%XMM17 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %RAX,0x1a8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMULSD 0x3a0(%RSP,%RAX,8),%XMM4,%XMM4 | 1 | 0.50 | 0.50 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 4 | 0.50 |
VMOVUPD %XMM4,0x260(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
CMP $0x81,%R10D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0xd8(%RSP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xd0(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JB 436d10 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1870> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
TEST %R10D,%R10D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 4361e8 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xd48> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VBROADCASTSD %XMM27,%YMM19 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM26,%YMM23 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM25,%YMM24 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD 0x260(%RSP),%XMM11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VBROADCASTSD %XMM11,%YMM21 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM17,%YMM22 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM16,%YMM28 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV 0x80(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x258(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x78(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x250(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x68(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x248(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x70(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x240(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x60(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x238(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R8,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R13,0x210(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDX,0x208(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x170(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x200(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xc0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1f8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x158(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1f0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x160(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1e8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x168(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1e0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R9,0x108(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x1a0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x230(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x198(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x228(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x190(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x220(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x188(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x218(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x300(%RSP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xf8(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1d8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xf0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1d0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xe8(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1c8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xe0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1c0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RAX,0x100(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
XOR %EBX,%EBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD 0x600(%RSP),%YMM13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5e0(%RSP),%YMM14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5c0(%RSP),%YMM15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
JMP 4365b6 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1116> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
TEST %R10D,%R10D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 4361e8 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xd48> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JE 436f62 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1ac2> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VBROADCASTSD %XMM27,%YMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD %YMM4,0x120(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VBROADCASTSD %XMM26,%YMM21 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM25,%YMM22 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD 0x260(%RSP),%YMM23 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.33 |
VBROADCASTSD %XMM17,%YMM24 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM16,%YMM28 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0xd8(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xc0(%RSP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xd0(%RSP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xf8(%RSP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xf0(%RSP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xe8(%RSP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xe0(%RSP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVAPD %YMM31,%YMM13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RCX,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV 0x90(%RSP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R10,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VMOVUPD 0x340(%RSP),%YMM9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x380(%RSP),%YMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
MOV 0x1b8(%RSP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD 0x360(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVSD 0x180(%RSP),%XMM28 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD 0x178(%RSP),%XMM24 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x88(%RSP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVAPD %YMM13,%YMM31 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
JE 4361e8 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xd48> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 436f64 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1ac4> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R10,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VPBROADCASTQ %RCX,%YMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VPCMPNLEUQ 0xad945(%RIP),%YMM4,%K1 | |||||||||||||||
KORTESTB %K1,%K1 | 1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
JE 43731a <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1e7a> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x150(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0xc8(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL 0x1a8(%RSP),%RDX | 1 | 0 | 1 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 1 |
ADD 0x2f8(%RSP),%RDX | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
ADD %RDX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x1b0(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD (%RSI,%RCX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
MOV 0x2e8(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %RDX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VMOVAPD %YMM5,%YMM13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD (%RSI,%RCX,8),%YMM5{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x760(%RSP),%YMM23 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM23{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD 0x780(%RSP),%YMM22 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM5,%YMM22{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
MOV 0x2e0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %RDX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VMOVUPD (%RSI,%RCX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x7a0(%RSP),%YMM10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM10{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
MOV 0x2f0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %RCX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VMOVUPD (%RSI,%RDX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x7c0(%RSP),%YMM9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM9{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMULPD %YMM8,%YMM23,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM14,%YMM22,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM15,%YMM10,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM9,%YMM30,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD (%R8,%RAX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMULPD %YMM29,%YMM23,%YMM21 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD 0x7e0(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMULPD %YMM11,%YMM22,%YMM4 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM13,%YMM9,%YMM4 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM7,%YMM10,%YMM21 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VBROADCASTSD %XMM27,%YMM5 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD %YMM6,0x7e0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM19,%YMM5 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM5,(%R8,%RAX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x80(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RCX,%RAX,1),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RCX,8),%YMM5{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VFMADD231PD %YMM4,%YMM12,%YMM21 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD 0x800(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM5,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VBROADCASTSD %XMM26,%YMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD %YMM6,0x800(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM19,%YMM4 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM4,(%R8,%RCX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VBROADCASTSD %XMM25,%YMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV 0x78(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RDX,%RAX,1),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RDX,8),%YMM5{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x820(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM5,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM4,%YMM5 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM6,0x820(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM21,%YMM5 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM5,(%R8,%RDX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x68(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RSI,%RAX,1),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RSI,8),%YMM5{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x840(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM5,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VBROADCASTSD 0x260(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.33 |
VMOVUPD %YMM6,0x840(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM19,%YMM5 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM5,(%R8,%RSI,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VBROADCASTSD %XMM17,%YMM5 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV 0x70(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RSI,%RAX,1),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RSI,8),%YMM6{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM11,%YMM25 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM7,%YMM11 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD 0x860(%RSP),%YMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM6,%YMM7{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM5,%YMM6 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM7,0x860(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM7,%YMM21,%YMM6 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM6,(%R8,%RSI,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x60(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RSI,%RAX,1),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RSI,8),%YMM6{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM15,%YMM17 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM14,%YMM15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM8,%YMM14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD 0x880(%RSP),%YMM8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM6,%YMM8{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM23,0x760(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VMULPD 0x460(%RSP),%YMM23,%YMM6 | 1 | 0.50 | 0.50 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 4 | 0.50 |
VMOVUPD %YMM22,0x780(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD231PD %YMM31,%YMM22,%YMM6 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM10,0x7a0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD231PD 0x440(%RSP),%YMM10,%YMM6 | 1 | 0.50 | 0.50 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 4 | 0.50 |
VMOVUPD 0x420(%RSP),%YMM10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD %YMM9,0x7c0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD231PD %YMM12,%YMM9,%YMM6 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD 0x340(%RSP),%YMM9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VBROADCASTSD %XMM16,%YMM7 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD %YMM8,0x880(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM8,%YMM7,%YMM6 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM6,(%R8,%RSI,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x20(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD (%R9,%RAX,8),%YMM6{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x8a0(%RSP),%YMM8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM6,%YMM8{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM8,0x8a0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM8,%YMM19,%YMM4 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVAPD %YMM14,%YMM8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM15,%YMM14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM17,%YMM15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM4,(%R9,%RAX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VMOVUPD (%R9,%RCX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x8c0(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM6,0x8c0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM19,%YMM5 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM5,(%R9,%RCX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VMOVUPD (%R9,%RDX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x8e0(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM5{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM5,0x8e0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM5,%YMM7,%YMM21 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM21,(%R9,%RDX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0xd8(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD (%RCX,%RAX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x900(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM5{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM5,0x900(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM5,%YMM7,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVAPD %YMM13,%YMM5 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM11,%YMM7 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM25,%YMM11 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM19,(%RCX,%RAX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x30(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 4361e8 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xd48> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
Function | miniqmcreference::einspline_spo_ref |
Source file and lines | MultiBsplineRef.hpp:227-262 |
Module | exec |
nb instructions | 280 |
nb uops | 287 |
loop length | 1919 |
used x86 registers | 14 |
used mmx registers | 0 |
used xmm registers | 9 |
used ymm registers | 23 |
used zmm registers | 0 |
nb stack references | 89 |
micro-operation queue | 48.00 cycles |
front end | 48.00 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 16.50 | 16.00 | 39.33 | 39.33 | 28.50 | 17.00 | 14.70 | 28.50 | 28.50 | 28.50 | 14.80 | 39.33 |
cycles | 16.50 | 16.00 | 39.33 | 39.33 | 28.50 | 17.00 | 14.70 | 28.50 | 28.50 | 28.50 | 14.80 | 39.33 |
Cycles executing div or sqrt instructions | NA |
FE+BE cycles | 49.65 |
Stall cycles | 0.00 |
Front-end | 48.00 |
Dispatch | 39.33 |
Overall L1 | 48.00 |
all | 2% |
load | 5% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 11% |
all | 80% |
load | 82% |
store | 100% |
mul | 40% |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | 100% |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 62% |
all | 60% |
load | 63% |
store | 45% |
mul | 40% |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | 100% |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 54% |
all | 12% |
load | 14% |
store | 12% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
other | 13% |
all | 42% |
load | 42% |
store | 49% |
mul | 27% |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | 50% |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 35% |
all | 35% |
load | 35% |
store | 29% |
mul | 27% |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | 50% |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 32% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
MOV 0x30(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD 0x620(%RSP),%YMM8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x600(%RSP),%YMM14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5e0(%RSP),%YMM15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5c0(%RSP),%YMM30 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5a0(%RSP),%YMM29 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x580(%RSP),%YMM11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x380(%RSP),%YMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x360(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVSD 0x180(%RSP),%XMM28 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD 0x178(%RSP),%XMM24 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x90(%RSP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x18(%RSP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x1a8(%RSP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA 0x1(%RDI),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x308(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RDX,0xe0(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0xe8(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0xf0(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0xf8(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
MOV 0xc8(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RDX,0x188(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0x190(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0x198(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
ADD %RDX,0x1a0(%RSP) | 2 | 0.20 | 0.20 | 0.33 | 0.33 | 0.50 | 0.20 | 0.20 | 0.50 | 0.50 | 0.50 | 0.20 | 0.33 | 1 | 0.50 |
CMP $0x3,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JE 436030 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xb90> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VMOVSD 0x540(%RSP,%RAX,8),%XMM16 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMULSD %XMM28,%XMM16,%XMM27 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMULSD %XMM24,%XMM16,%XMM25 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVSD 0x310(%RSP),%XMM4 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMULSD %XMM4,%XMM16,%XMM16 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVSD 0x3e0(%RSP,%RAX,8),%XMM17 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMULSD %XMM24,%XMM17,%XMM26 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMULSD %XMM4,%XMM17,%XMM17 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
MOV %RAX,0x1a8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
VMULSD 0x3a0(%RSP,%RAX,8),%XMM4,%XMM4 | 1 | 0.50 | 0.50 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 4 | 0.50 |
VMOVUPD %XMM4,0x260(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
CMP $0x81,%R10D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0xd8(%RSP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xd0(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JB 436d10 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1870> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
TEST %R10D,%R10D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 4361e8 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xd48> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VBROADCASTSD %XMM27,%YMM19 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM26,%YMM23 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM25,%YMM24 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD 0x260(%RSP),%XMM11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VBROADCASTSD %XMM11,%YMM21 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM17,%YMM22 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM16,%YMM28 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV 0x80(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x258(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x78(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x250(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x68(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x248(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x70(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x240(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x60(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x238(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R8,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R13,0x210(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDX,0x208(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x170(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x200(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xc0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1f8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x158(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1f0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x160(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1e8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x168(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1e0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R9,0x108(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x1a0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x230(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x198(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x228(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x190(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x220(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x188(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x218(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0x300(%RSP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xf8(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1d8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xf0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1d0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xe8(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1c8(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV 0xe0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RCX,0x1c0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RAX,0x100(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
XOR %EBX,%EBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD 0x600(%RSP),%YMM13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5e0(%RSP),%YMM14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x5c0(%RSP),%YMM15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
JMP 4365b6 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1116> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
TEST %R10D,%R10D | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JLE 4361e8 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xd48> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JE 436f62 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1ac2> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
VBROADCASTSD %XMM27,%YMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD %YMM4,0x120(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VBROADCASTSD %XMM26,%YMM21 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM25,%YMM22 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD 0x260(%RSP),%YMM23 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.33 |
VBROADCASTSD %XMM17,%YMM24 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VBROADCASTSD %XMM16,%YMM28 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0xd8(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xc0(%RSP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xd0(%RSP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xf8(%RSP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xf0(%RSP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xe8(%RSP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0xe0(%RSP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVAPD %YMM31,%YMM13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RCX,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV 0x90(%RSP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R10,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VMOVUPD 0x340(%RSP),%YMM9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x380(%RSP),%YMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
MOV 0x1b8(%RSP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD 0x360(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVSD 0x180(%RSP),%XMM28 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVSD 0x178(%RSP),%XMM24 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x88(%RSP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVAPD %YMM13,%YMM31 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
JE 4361e8 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xd48> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 436f64 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1ac4> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R10,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VPBROADCASTQ %RCX,%YMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VPCMPNLEUQ 0xad945(%RIP),%YMM4,%K1 | |||||||||||||||
KORTESTB %K1,%K1 | 1 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 |
JE 43731a <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0x1e7a> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x150(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0xc8(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL 0x1a8(%RSP),%RDX | 1 | 0 | 1 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 1 |
ADD 0x2f8(%RSP),%RDX | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
ADD %RDX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x1b0(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD (%RSI,%RCX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
MOV 0x2e8(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %RDX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VMOVAPD %YMM5,%YMM13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD (%RSI,%RCX,8),%YMM5{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x760(%RSP),%YMM23 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM23{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD 0x780(%RSP),%YMM22 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM5,%YMM22{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
MOV 0x2e0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %RDX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VMOVUPD (%RSI,%RCX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x7a0(%RSP),%YMM10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM10{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
MOV 0x2f0(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
ADD %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %RCX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
VMOVUPD (%RSI,%RDX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x7c0(%RSP),%YMM9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM9{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMULPD %YMM8,%YMM23,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM14,%YMM22,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM15,%YMM10,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM9,%YMM30,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD (%R8,%RAX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMULPD %YMM29,%YMM23,%YMM21 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD 0x7e0(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMULPD %YMM11,%YMM22,%YMM4 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM13,%YMM9,%YMM4 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VFMADD231PD %YMM7,%YMM10,%YMM21 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VBROADCASTSD %XMM27,%YMM5 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD %YMM6,0x7e0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM19,%YMM5 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM5,(%R8,%RAX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x80(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RCX,%RAX,1),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RCX,8),%YMM5{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VFMADD231PD %YMM4,%YMM12,%YMM21 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD 0x800(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM5,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VBROADCASTSD %XMM26,%YMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD %YMM6,0x800(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM19,%YMM4 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM4,(%R8,%RCX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VBROADCASTSD %XMM25,%YMM4 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV 0x78(%RSP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RDX,%RAX,1),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RDX,8),%YMM5{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x820(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM5,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM4,%YMM5 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM6,0x820(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM21,%YMM5 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM5,(%R8,%RDX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x68(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RSI,%RAX,1),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RSI,8),%YMM5{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x840(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM5,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VBROADCASTSD 0x260(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 3 | 0.33 |
VMOVUPD %YMM6,0x840(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM19,%YMM5 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM5,(%R8,%RSI,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VBROADCASTSD %XMM17,%YMM5 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
MOV 0x70(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RSI,%RAX,1),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RSI,8),%YMM6{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM11,%YMM25 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM7,%YMM11 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD 0x860(%RSP),%YMM7 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM6,%YMM7{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM5,%YMM6 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM7,0x860(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM7,%YMM21,%YMM6 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM6,(%R8,%RSI,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x60(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (%RSI,%RAX,1),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
VMOVUPD (%R8,%RSI,8),%YMM6{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM15,%YMM17 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM14,%YMM15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM8,%YMM14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD 0x880(%RSP),%YMM8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM6,%YMM8{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM23,0x760(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VMULPD 0x460(%RSP),%YMM23,%YMM6 | 1 | 0.50 | 0.50 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 4 | 0.50 |
VMOVUPD %YMM22,0x780(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD231PD %YMM31,%YMM22,%YMM6 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM10,0x7a0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD231PD 0x440(%RSP),%YMM10,%YMM6 | 1 | 0.50 | 0.50 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 4 | 0.50 |
VMOVUPD 0x420(%RSP),%YMM10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD %YMM9,0x7c0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD231PD %YMM12,%YMM9,%YMM6 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD 0x340(%RSP),%YMM9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VBROADCASTSD %XMM16,%YMM7 | 1 | 0 | 0 | 0 | 0 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
VMOVUPD %YMM8,0x880(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM8,%YMM7,%YMM6 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM6,(%R8,%RSI,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x20(%RSP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD (%R9,%RAX,8),%YMM6{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x8a0(%RSP),%YMM8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM6,%YMM8{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM8,0x8a0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM8,%YMM19,%YMM4 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVAPD %YMM14,%YMM8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM15,%YMM14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM17,%YMM15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM4,(%R9,%RAX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VMOVUPD (%R9,%RCX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x8c0(%RSP),%YMM6 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM6{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM6,0x8c0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM6,%YMM19,%YMM5 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM5,(%R9,%RCX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VMOVUPD (%R9,%RDX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x8e0(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM5{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM5,0x8e0(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM5,%YMM7,%YMM21 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVUPD %YMM21,(%R9,%RDX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0xd8(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
VMOVUPD (%RCX,%RAX,8),%YMM4{%K1}{z} | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVUPD 0x900(%RSP),%YMM5 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0-1 | 0.33 |
VMOVAPD %YMM4,%YMM5{%K1} | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM5,0x900(%RSP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
VFMADD213PD %YMM5,%YMM7,%YMM19 | 1 | 0.50 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 |
VMOVAPD %YMM13,%YMM5 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM11,%YMM7 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVAPD %YMM25,%YMM11 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0-1 | 0.17 |
VMOVUPD %YMM19,(%RCX,%RAX,8){%K1} | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0-1 | 0.50 |
MOV 0x30(%RSP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JMP 4361e8 <_ZN16miniqmcreference17einspline_spo_refIdE8evaluateERKN11qmcplusplus11ParticleSetEiRNS2_6VectorIdSaIdEEERNS6_INS2_10TinyVectorIdLj3EEESaISB_EEES9_+0xd48> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |