Loop Id: 851 | Module: libparcsr_ls.so | Source: par_multi_interp.c:1072-1125 | Coverage: 0.08% |
---|
Loop Id: 851 | Module: libparcsr_ls.so | Source: par_multi_interp.c:1072-1125 | Coverage: 0.08% |
---|
0x7fc50 MOV -0xb0(%RBP),%R13 |
0x7fc57 MOV -0x70(%RBP),%RCX |
0x7fc5b MOV -0xd0(%RBP),%R8 |
0x7fc62 MOV (%R13,%RCX,8),%RDI |
0x7fc67 LEA 0x8(,%RDI,8),%RDX |
0x7fc6f MOV (%R8,%RDI,8),%R11 |
0x7fc73 MOV %RDI,%RAX |
0x7fc76 ADD %RDX,%R8 |
0x7fc79 NOT %RAX |
0x7fc7c MOV %R8,-0x90(%RBP) |
0x7fc83 MOV (%R8),%R8 |
0x7fc86 MOV %R11,-0x38(%RBP) |
0x7fc8a CMP %R8,%R11 |
0x7fc8d JGE 7ffb0 |
0x7fc93 MOV %RDI,-0xe0(%RBP) |
0x7fc9a MOV %RBX,%R13 |
0x7fc9d MOV -0x88(%RBP),%RDI |
0x7fca4 MOV %RDX,-0xc0(%RBP) |
0x7fcab JMP 7fcc1 |
(854) 0x7fcb0 INCQ -0x38(%RBP) |
(854) 0x7fcb4 MOV -0x38(%RBP),%RDX |
(854) 0x7fcb8 CMP %R8,%RDX |
(854) 0x7fcbb JGE 7ff98 |
(854) 0x7fcc1 MOV -0x40(%RBP),%RDX |
(854) 0x7fcc5 MOV -0x38(%RBP),%RBX |
(854) 0x7fcc9 MOV (%RDX,%RBX,8),%RCX |
(854) 0x7fccd MOV -0x48(%RBP),%RDX |
(854) 0x7fcd1 LEA (,%RCX,8),%R11 |
(854) 0x7fcd9 CMP %R13,(%RDX,%RCX,8) |
(854) 0x7fcdd JNE 7fcb0 |
(854) 0x7fcdf MOV -0x80(%RBP),%R8 |
(854) 0x7fce3 MOV -0x60(%RBP),%RBX |
(854) 0x7fce7 MOV (%R8,%RCX,8),%R8 |
(854) 0x7fceb MOV 0x8(%RBX,%R11,1),%RBX |
(854) 0x7fcf0 ADD %R8,%RBX |
(854) 0x7fcf3 CMP %RBX,%R8 |
(854) 0x7fcf6 JGE 7fe54 |
(854) 0x7fcfc MOV -0x8(%R10,%R14,1),%RDX |
(854) 0x7fd01 LEA (%RDX,%RBX,8),%RBX |
(854) 0x7fd05 LEA (%RDX,%R8,8),%R8 |
(854) 0x7fd09 MOV %RBX,-0x68(%RBP) |
(854) 0x7fd0d SUB %R8,%RBX |
(854) 0x7fd10 SUB $0x8,%RBX |
(854) 0x7fd14 SHR $0x3,%RBX |
(854) 0x7fd18 INC %RBX |
(854) 0x7fd1b AND $0x3,%EBX |
(854) 0x7fd1e JE 7fdc7 |
(854) 0x7fd24 CMP $0x1,%RBX |
(854) 0x7fd28 JE 7fd8e |
(854) 0x7fd2a CMP $0x2,%RBX |
(854) 0x7fd2e JE 7fd5f |
(854) 0x7fd30 MOV (%R8),%RDX |
(854) 0x7fd33 MOV %RDX,%RBX |
(854) 0x7fd36 LEA (%R12,%RDX,8),%RDX |
(854) 0x7fd3a MOV %RDX,-0x88(%RBP) |
(854) 0x7fd41 CMP %RAX,(%RDX) |
(854) 0x7fd44 JE 7fd5b |
(854) 0x7fd46 MOV (%R10,%R14,1),%RDX |
(854) 0x7fd4a MOV %RBX,(%RDX,%RSI,8) |
(854) 0x7fd4e MOV -0x88(%RBP),%RBX |
(854) 0x7fd55 INC %RSI |
(854) 0x7fd58 MOV %RAX,(%RBX) |
(854) 0x7fd5b ADD $0x8,%R8 |
(854) 0x7fd5f MOV (%R8),%RDX |
(854) 0x7fd62 MOV %RDX,%RBX |
(854) 0x7fd65 LEA (%R12,%RDX,8),%RDX |
(854) 0x7fd69 MOV %RDX,-0x88(%RBP) |
(854) 0x7fd70 CMP %RAX,(%RDX) |
(854) 0x7fd73 JE 7fd8a |
(854) 0x7fd75 MOV (%R10,%R14,1),%RDX |
(854) 0x7fd79 MOV %RBX,(%RDX,%RSI,8) |
(854) 0x7fd7d MOV -0x88(%RBP),%RBX |
(854) 0x7fd84 INC %RSI |
(854) 0x7fd87 MOV %RAX,(%RBX) |
(854) 0x7fd8a ADD $0x8,%R8 |
(854) 0x7fd8e MOV (%R8),%RDX |
(854) 0x7fd91 MOV %RDX,%RBX |
(854) 0x7fd94 LEA (%R12,%RDX,8),%RDX |
(854) 0x7fd98 MOV %RDX,-0x88(%RBP) |
(854) 0x7fd9f CMP %RAX,(%RDX) |
(854) 0x7fda2 JE 7fdb9 |
(854) 0x7fda4 MOV (%R10,%R14,1),%RDX |
(854) 0x7fda8 MOV %RBX,(%RDX,%RSI,8) |
(854) 0x7fdac MOV -0x88(%RBP),%RBX |
(854) 0x7fdb3 INC %RSI |
(854) 0x7fdb6 MOV %RAX,(%RBX) |
(854) 0x7fdb9 ADD $0x8,%R8 |
(854) 0x7fdbd CMP %R8,-0x68(%RBP) |
(854) 0x7fdc1 JE 7fe54 |
(854) 0x7fdc7 MOV %RCX,-0x88(%RBP) |
(854) 0x7fdce MOV %RDI,%RBX |
(856) 0x7fdd1 MOV (%R8),%RDI |
(856) 0x7fdd4 LEA (%R12,%RDI,8),%RDX |
(856) 0x7fdd8 CMP %RAX,(%RDX) |
(856) 0x7fddb JE 7fdeb |
(856) 0x7fddd MOV (%R10,%R14,1),%RCX |
(856) 0x7fde1 MOV %RDI,(%RCX,%RSI,8) |
(856) 0x7fde5 INC %RSI |
(856) 0x7fde8 MOV %RAX,(%RDX) |
(856) 0x7fdeb LEA 0x8(%R8),%RDX |
(856) 0x7fdef MOV 0x8(%R8),%R8 |
(856) 0x7fdf3 LEA (%R12,%R8,8),%RDI |
(856) 0x7fdf7 CMP %RAX,(%RDI) |
(856) 0x7fdfa JE 7fe0a |
(856) 0x7fdfc MOV (%R10,%R14,1),%RCX |
(856) 0x7fe00 MOV %R8,(%RCX,%RSI,8) |
(856) 0x7fe04 INC %RSI |
(856) 0x7fe07 MOV %RAX,(%RDI) |
(856) 0x7fe0a MOV 0x8(%RDX),%R8 |
(856) 0x7fe0e LEA (%R12,%R8,8),%RDI |
(856) 0x7fe12 CMP %RAX,(%RDI) |
(856) 0x7fe15 JE 7fe25 |
(856) 0x7fe17 MOV (%R10,%R14,1),%RCX |
(856) 0x7fe1b MOV %R8,(%RCX,%RSI,8) |
(856) 0x7fe1f INC %RSI |
(856) 0x7fe22 MOV %RAX,(%RDI) |
(856) 0x7fe25 MOV 0x10(%RDX),%R8 |
(856) 0x7fe29 LEA (%R12,%R8,8),%RDI |
(856) 0x7fe2d CMP %RAX,(%RDI) |
(856) 0x7fe30 JE 7fe40 |
(856) 0x7fe32 MOV (%R10,%R14,1),%RCX |
(856) 0x7fe36 MOV %R8,(%RCX,%RSI,8) |
(856) 0x7fe3a INC %RSI |
(856) 0x7fe3d MOV %RAX,(%RDI) |
(856) 0x7fe40 LEA 0x18(%RDX),%R8 |
(856) 0x7fe44 CMP %R8,-0x68(%RBP) |
(856) 0x7fe48 JNE 7fdd1 |
(854) 0x7fe4a MOV -0x88(%RBP),%RCX |
(854) 0x7fe51 MOV %RBX,%RDI |
(854) 0x7fe54 MOV -0x78(%RBP),%RBX |
(854) 0x7fe58 MOV -0x58(%RBP),%RDX |
(854) 0x7fe5c MOV (%RBX,%RCX,8),%RCX |
(854) 0x7fe60 MOV 0x8(%RDX,%R11,1),%R8 |
(854) 0x7fe65 ADD %RCX,%R8 |
(854) 0x7fe68 CMP %R8,%RCX |
(854) 0x7fe6b JGE 80390 |
(854) 0x7fe71 MOV -0x8(%R9,%R14,1),%RBX |
(854) 0x7fe76 LEA (%RBX,%RCX,8),%RCX |
(854) 0x7fe7a LEA (%RBX,%R8,8),%RBX |
(854) 0x7fe7e MOV %RBX,%RDX |
(854) 0x7fe81 SUB %RCX,%RDX |
(854) 0x7fe84 SUB $0x8,%RDX |
(854) 0x7fe88 SHR $0x3,%RDX |
(854) 0x7fe8c INC %RDX |
(854) 0x7fe8f AND $0x3,%EDX |
(854) 0x7fe92 JE 7feff |
(854) 0x7fe94 CMP $0x1,%RDX |
(854) 0x7fe98 JE 7fedc |
(854) 0x7fe9a CMP $0x2,%RDX |
(854) 0x7fe9e JE 7febe |
(854) 0x7fea0 MOV (%RCX),%R11 |
(854) 0x7fea3 LEA (%R15,%R11,8),%R8 |
(854) 0x7fea7 CMP %RAX,(%R8) |
(854) 0x7feaa JE 7feba |
(854) 0x7feac MOV (%R9,%R14,1),%RDX |
(854) 0x7feb0 MOV %R11,(%RDX,%RDI,8) |
(854) 0x7feb4 INC %RDI |
(854) 0x7feb7 MOV %RAX,(%R8) |
(854) 0x7feba ADD $0x8,%RCX |
(854) 0x7febe MOV (%RCX),%R11 |
(854) 0x7fec1 LEA (%R15,%R11,8),%R8 |
(854) 0x7fec5 CMP %RAX,(%R8) |
(854) 0x7fec8 JE 7fed8 |
(854) 0x7feca MOV (%R9,%R14,1),%RDX |
(854) 0x7fece MOV %R11,(%RDX,%RDI,8) |
(854) 0x7fed2 INC %RDI |
(854) 0x7fed5 MOV %RAX,(%R8) |
(854) 0x7fed8 ADD $0x8,%RCX |
(854) 0x7fedc MOV (%RCX),%R11 |
(854) 0x7fedf LEA (%R15,%R11,8),%R8 |
(854) 0x7fee3 CMP %RAX,(%R8) |
(854) 0x7fee6 JE 7fef6 |
(854) 0x7fee8 MOV (%R9,%R14,1),%RDX |
(854) 0x7feec MOV %R11,(%RDX,%RDI,8) |
(854) 0x7fef0 INC %RDI |
(854) 0x7fef3 MOV %RAX,(%R8) |
(854) 0x7fef6 ADD $0x8,%RCX |
(854) 0x7fefa CMP %RBX,%RCX |
(854) 0x7fefd JE 7ff77 |
(855) 0x7feff MOV (%RCX),%R8 |
(855) 0x7ff02 LEA (%R15,%R8,8),%RDX |
(855) 0x7ff06 CMP %RAX,(%RDX) |
(855) 0x7ff09 JE 7ff19 |
(855) 0x7ff0b MOV (%R9,%R14,1),%R11 |
(855) 0x7ff0f MOV %R8,(%R11,%RDI,8) |
(855) 0x7ff13 INC %RDI |
(855) 0x7ff16 MOV %RAX,(%RDX) |
(855) 0x7ff19 MOV 0x8(%RCX),%R8 |
(855) 0x7ff1d LEA 0x8(%RCX),%RDX |
(855) 0x7ff21 LEA (%R15,%R8,8),%RCX |
(855) 0x7ff25 CMP %RAX,(%RCX) |
(855) 0x7ff28 JE 7ff38 |
(855) 0x7ff2a MOV (%R9,%R14,1),%R11 |
(855) 0x7ff2e MOV %R8,(%R11,%RDI,8) |
(855) 0x7ff32 INC %RDI |
(855) 0x7ff35 MOV %RAX,(%RCX) |
(855) 0x7ff38 MOV 0x8(%RDX),%R8 |
(855) 0x7ff3c LEA (%R15,%R8,8),%RCX |
(855) 0x7ff40 CMP %RAX,(%RCX) |
(855) 0x7ff43 JE 7ff53 |
(855) 0x7ff45 MOV (%R9,%R14,1),%R11 |
(855) 0x7ff49 MOV %R8,(%R11,%RDI,8) |
(855) 0x7ff4d INC %RDI |
(855) 0x7ff50 MOV %RAX,(%RCX) |
(855) 0x7ff53 MOV 0x10(%RDX),%R8 |
(855) 0x7ff57 LEA (%R15,%R8,8),%RCX |
(855) 0x7ff5b CMP %RAX,(%RCX) |
(855) 0x7ff5e JE 7ff6e |
(855) 0x7ff60 MOV (%R9,%R14,1),%R11 |
(855) 0x7ff64 MOV %R8,(%R11,%RDI,8) |
(855) 0x7ff68 INC %RDI |
(855) 0x7ff6b MOV %RAX,(%RCX) |
(855) 0x7ff6e LEA 0x18(%RDX),%RCX |
(855) 0x7ff72 CMP %RBX,%RCX |
(855) 0x7ff75 JNE 7feff |
(854) 0x7ff77 MOV -0x90(%RBP),%RBX |
(854) 0x7ff7e INCQ -0x38(%RBP) |
(854) 0x7ff82 MOV -0x38(%RBP),%RDX |
(854) 0x7ff86 MOV (%RBX),%R8 |
(854) 0x7ff89 CMP %R8,%RDX |
(854) 0x7ff8c JL 7fcc1 |
0x7ff92 NOPW (%RAX,%RAX,1) |
0x7ff98 MOV %RDI,-0x88(%RBP) |
0x7ff9f MOV -0xc0(%RBP),%RDX |
0x7ffa6 MOV %R13,%RBX |
0x7ffa9 MOV -0xe0(%RBP),%RDI |
0x7ffb0 MOV -0xc8(%RBP),%R11 |
0x7ffb7 MOV (%R11,%RDI,8),%RAX |
0x7ffbb ADD %RDX,%R11 |
0x7ffbe NOT %RDI |
0x7ffc1 MOV (%R11),%RCX |
0x7ffc4 MOV %R11,-0xc0(%RBP) |
0x7ffcb CMP %RAX,%RCX |
0x7ffce JLE 8028a |
0x7ffd4 MOV %R9,-0x68(%RBP) |
0x7ffd8 MOV -0x88(%RBP),%R8 |
0x7ffdf MOV %R10,-0x90(%RBP) |
0x7ffe6 MOV -0xd8(%RBP),%R13 |
0x7ffed JMP 7fffc |
(852) 0x7fff0 INC %RAX |
(852) 0x7fff3 CMP %RCX,%RAX |
(852) 0x7fff6 JGE 80278 |
(852) 0x7fffc MOV -0x50(%RBP),%R9 |
(852) 0x80000 MOV (%R9,%RAX,8),%RDX |
(852) 0x80004 LEA (,%RDX,8),%R11 |
(852) 0x8000c CMP %RBX,(%R13,%RDX,8) |
(852) 0x80011 JNE 7fff0 |
(852) 0x80013 MOV -0xa0(%RBP),%R10 |
(852) 0x8001a MOV -0x98(%RBP),%R9 |
(852) 0x80021 MOV (%R10,%RDX,8),%RDX |
(852) 0x80025 MOV 0x8(%R9,%R11,1),%R11 |
(852) 0x8002a ADD %RDX,%R11 |
(852) 0x8002d CMP %R11,%RDX |
(852) 0x80030 JGE 7fff0 |
(852) 0x80032 MOV -0xa8(%RBP),%RCX |
(852) 0x80039 MOV (%RCX,%R14,1),%R10 |
(852) 0x8003d LEA (%R10,%R11,8),%R9 |
(852) 0x80041 LEA (%R10,%RDX,8),%RDX |
(852) 0x80045 MOV %R9,-0x38(%RBP) |
(852) 0x80049 SUB %RDX,%R9 |
(852) 0x8004c SUB $0x8,%R9 |
(852) 0x80050 SHR $0x3,%R9 |
(852) 0x80054 INC %R9 |
(852) 0x80057 AND $0x3,%R9D |
(852) 0x8005b JE 800fb |
(852) 0x80061 CMP $0x1,%R9 |
(852) 0x80065 JE 800c3 |
(852) 0x80067 CMP $0x2,%R9 |
(852) 0x8006b JE 80098 |
(852) 0x8006d MOV (%RDX),%RCX |
(852) 0x80070 TEST %RCX,%RCX |
(852) 0x80073 JS 803a0 |
(852) 0x80079 LEA (%R15,%RCX,8),%R11 |
(852) 0x8007d CMP %RDI,(%R11) |
(852) 0x80080 JE 80094 |
(852) 0x80082 MOV -0x68(%RBP),%R10 |
(852) 0x80086 MOV (%R10,%R14,1),%R9 |
(852) 0x8008a MOV %RCX,(%R9,%R8,8) |
(852) 0x8008e INC %R8 |
(852) 0x80091 MOV %RDI,(%R11) |
(852) 0x80094 ADD $0x8,%RDX |
(852) 0x80098 MOV (%RDX),%RCX |
(852) 0x8009b TEST %RCX,%RCX |
(852) 0x8009e JS 80360 |
(852) 0x800a4 LEA (%R15,%RCX,8),%R10 |
(852) 0x800a8 CMP %RDI,(%R10) |
(852) 0x800ab JE 800bf |
(852) 0x800ad MOV -0x68(%RBP),%R9 |
(852) 0x800b1 MOV (%R9,%R14,1),%R11 |
(852) 0x800b5 MOV %RCX,(%R11,%R8,8) |
(852) 0x800b9 INC %R8 |
(852) 0x800bc MOV %RDI,(%R10) |
(852) 0x800bf ADD $0x8,%RDX |
(852) 0x800c3 MOV (%RDX),%RCX |
(852) 0x800c6 TEST %RCX,%RCX |
(852) 0x800c9 JS 80328 |
(852) 0x800cf LEA (%R15,%RCX,8),%R11 |
(852) 0x800d3 CMP %RDI,(%R11) |
(852) 0x800d6 JE 800ea |
(852) 0x800d8 MOV -0x68(%RBP),%R9 |
(852) 0x800dc MOV (%R9,%R14,1),%R10 |
(852) 0x800e0 MOV %RCX,(%R10,%R8,8) |
(852) 0x800e4 INC %R8 |
(852) 0x800e7 MOV %RDI,(%R11) |
(852) 0x800ea MOV -0x38(%RBP),%RCX |
(852) 0x800ee ADD $0x8,%RDX |
(852) 0x800f2 CMP %RCX,%RDX |
(852) 0x800f5 JE 8025d |
(852) 0x800fb MOV %RAX,-0x88(%RBP) |
(852) 0x80102 MOV -0x68(%RBP),%R9 |
(852) 0x80106 MOV -0x90(%RBP),%R10 |
(852) 0x8010d JMP 80180 |
(853) 0x80110 LEA (%R15,%RDX,8),%R11 |
(853) 0x80114 CMP %RDI,(%R11) |
(853) 0x80117 JE 80127 |
(853) 0x80119 MOV (%R9,%R14,1),%RAX |
(853) 0x8011d MOV %RDX,(%RAX,%R8,8) |
(853) 0x80121 INC %R8 |
(853) 0x80124 MOV %RDI,(%R11) |
(853) 0x80127 MOV 0x8(%RCX),%RDX |
(853) 0x8012b TEST %RDX,%RDX |
(853) 0x8012e JS 801e5 |
(853) 0x80134 LEA (%R15,%RDX,8),%R11 |
(853) 0x80138 CMP %RDI,(%R11) |
(853) 0x8013b JE 8014b |
(853) 0x8013d MOV (%R9,%R14,1),%RAX |
(853) 0x80141 MOV %RDX,(%RAX,%R8,8) |
(853) 0x80145 INC %R8 |
(853) 0x80148 MOV %RDI,(%R11) |
(853) 0x8014b MOV 0x10(%RCX),%RDX |
(853) 0x8014f TEST %RDX,%RDX |
(853) 0x80152 JS 80216 |
(853) 0x80158 LEA (%R15,%RDX,8),%R11 |
(853) 0x8015c CMP %RDI,(%R11) |
(853) 0x8015f JE 8016f |
(853) 0x80161 MOV (%R9,%R14,1),%RAX |
(853) 0x80165 MOV %RDX,(%RAX,%R8,8) |
(853) 0x80169 INC %R8 |
(853) 0x8016c MOV %RDI,(%R11) |
(853) 0x8016f LEA 0x18(%RCX),%RDX |
(853) 0x80173 MOV -0x38(%RBP),%RCX |
(853) 0x80177 CMP %RCX,%RDX |
(853) 0x8017a JE 8024b |
(853) 0x80180 MOV (%RDX),%RCX |
(853) 0x80183 TEST %RCX,%RCX |
(853) 0x80186 JS 802f8 |
(853) 0x8018c LEA (%R15,%RCX,8),%R11 |
(853) 0x80190 CMP %RDI,(%R11) |
(853) 0x80193 JE 801a3 |
(853) 0x80195 MOV (%R9,%R14,1),%RAX |
(853) 0x80199 MOV %RCX,(%RAX,%R8,8) |
(853) 0x8019d INC %R8 |
(853) 0x801a0 MOV %RDI,(%R11) |
(853) 0x801a3 LEA 0x8(%RDX),%RCX |
(853) 0x801a7 MOV 0x8(%RDX),%RDX |
(853) 0x801ab TEST %RDX,%RDX |
(853) 0x801ae JNS 80110 |
(853) 0x801b4 MOV %RDX,%R11 |
(853) 0x801b7 NOT %R11 |
(853) 0x801ba LEA (%R12,%R11,8),%R11 |
(853) 0x801be CMP %RDI,(%R11) |
(853) 0x801c1 JE 80127 |
(853) 0x801c7 MOV (%R10,%R14,1),%RAX |
(853) 0x801cb NOT %RDX |
(853) 0x801ce MOV %RDX,(%RAX,%RSI,8) |
(853) 0x801d2 INC %RSI |
(853) 0x801d5 MOV %RDI,(%R11) |
(853) 0x801d8 MOV 0x8(%RCX),%RDX |
(853) 0x801dc TEST %RDX,%RDX |
(853) 0x801df JNS 80134 |
(853) 0x801e5 MOV %RDX,%R11 |
(853) 0x801e8 NOT %R11 |
(853) 0x801eb LEA (%R12,%R11,8),%R11 |
(853) 0x801ef CMP %RDI,(%R11) |
(853) 0x801f2 JE 8014b |
(853) 0x801f8 MOV (%R10,%R14,1),%RAX |
(853) 0x801fc NOT %RDX |
(853) 0x801ff MOV %RDX,(%RAX,%RSI,8) |
(853) 0x80203 INC %RSI |
(853) 0x80206 MOV %RDI,(%R11) |
(853) 0x80209 MOV 0x10(%RCX),%RDX |
(853) 0x8020d TEST %RDX,%RDX |
(853) 0x80210 JNS 80158 |
(853) 0x80216 MOV %RDX,%R11 |
(853) 0x80219 NOT %R11 |
(853) 0x8021c LEA (%R12,%R11,8),%R11 |
(853) 0x80220 CMP %RDI,(%R11) |
(853) 0x80223 JE 8016f |
(853) 0x80229 MOV (%R10,%R14,1),%RAX |
(853) 0x8022d NOT %RDX |
(853) 0x80230 MOV %RDX,(%RAX,%RSI,8) |
(853) 0x80234 LEA 0x18(%RCX),%RDX |
(853) 0x80238 MOV -0x38(%RBP),%RCX |
(853) 0x8023c INC %RSI |
(853) 0x8023f MOV %RDI,(%R11) |
(853) 0x80242 CMP %RCX,%RDX |
(853) 0x80245 JNE 80180 |
(852) 0x8024b MOV %R9,-0x68(%RBP) |
(852) 0x8024f MOV -0x88(%RBP),%RAX |
(852) 0x80256 MOV %R10,-0x90(%RBP) |
(852) 0x8025d MOV -0xc0(%RBP),%R9 |
(852) 0x80264 INC %RAX |
(852) 0x80267 MOV (%R9),%RCX |
(852) 0x8026a CMP %RCX,%RAX |
(852) 0x8026d JL 7fffc |
0x80273 NOPL (%RAX,%RAX,1) |
0x80278 MOV %R8,-0x88(%RBP) |
0x8027f MOV -0x68(%RBP),%R9 |
0x80283 MOV -0x90(%RBP),%R10 |
0x8028a INCQ -0x70(%RBP) |
0x8028e MOV -0x70(%RBP),%RDI |
0x80292 CMP %RDI,-0xb8(%RBP) |
0x80299 JG 7fc50 |
(853) 0x802f8 MOV %RCX,%R11 |
(853) 0x802fb NOT %R11 |
(853) 0x802fe LEA (%R12,%R11,8),%R11 |
(853) 0x80302 CMP %RDI,(%R11) |
(853) 0x80305 JE 801a3 |
(853) 0x8030b MOV (%R10,%R14,1),%RAX |
(853) 0x8030f NOT %RCX |
(853) 0x80312 MOV %RCX,(%RAX,%RSI,8) |
(853) 0x80316 INC %RSI |
(853) 0x80319 MOV %RDI,(%R11) |
(853) 0x8031c JMP 801a3 |
(852) 0x80328 MOV %RCX,%R11 |
(852) 0x8032b NOT %R11 |
(852) 0x8032e LEA (%R12,%R11,8),%R10 |
(852) 0x80332 CMP %RDI,(%R10) |
(852) 0x80335 JE 800ea |
(852) 0x8033b MOV -0x90(%RBP),%R9 |
(852) 0x80342 NOT %RCX |
(852) 0x80345 MOV (%R9,%R14,1),%R11 |
(852) 0x80349 MOV %RCX,(%R11,%RSI,8) |
(852) 0x8034d INC %RSI |
(852) 0x80350 MOV %RDI,(%R10) |
(852) 0x80353 JMP 800ea |
(852) 0x80360 MOV %RCX,%R10 |
(852) 0x80363 NOT %R10 |
(852) 0x80366 LEA (%R12,%R10,8),%R11 |
(852) 0x8036a CMP %RDI,(%R11) |
(852) 0x8036d JE 800bf |
(852) 0x80373 MOV -0x90(%RBP),%R9 |
(852) 0x8037a NOT %RCX |
(852) 0x8037d MOV (%R9,%R14,1),%R10 |
(852) 0x80381 MOV %RCX,(%R10,%RSI,8) |
(852) 0x80385 INC %RSI |
(852) 0x80388 MOV %RDI,(%R11) |
(852) 0x8038b JMP 800bf |
(854) 0x80390 MOV -0x90(%RBP),%R11 |
(854) 0x80397 MOV (%R11),%R8 |
(854) 0x8039a JMP 7fcb0 |
(852) 0x803a0 MOV %RCX,%R11 |
(852) 0x803a3 NOT %R11 |
(852) 0x803a6 LEA (%R12,%R11,8),%R10 |
(852) 0x803aa CMP %RDI,(%R10) |
(852) 0x803ad JE 80094 |
(852) 0x803b3 MOV -0x90(%RBP),%R9 |
(852) 0x803ba NOT %RCX |
(852) 0x803bd MOV (%R9,%R14,1),%R11 |
(852) 0x803c1 MOV %RCX,(%R11,%RSI,8) |
(852) 0x803c5 INC %RSI |
(852) 0x803c8 MOV %RDI,(%R10) |
(852) 0x803cb JMP 80094 |
/home/eoseret/qaas_runs_CPU_9468/172-019-1763/intel/AMG/build/AMG/AMG/parcsr_ls/par_multi_interp.c: 1072 - 1125 |
-------------------------------------------------------------------------------- |
1072: for (i=thread_start; i < thread_stop; i++) |
1073: { |
1074: i1 = pass_array[i]; |
1075: for (j=S_diag_i[i1]; j < S_diag_i[i1+1]; j++) |
1076: { |
1077: j1 = S_diag_j[j]; |
1078: if (assigned[j1] == pass-1) |
1079: { |
1080: j_start = P_diag_start[j1]; |
1081: j_end = j_start+P_diag_i[j1+1]; |
1082: for (k=j_start; k < j_end; k++) |
1083: { |
1084: k1 = P_diag_pass[pass-1][k]; |
1085: if (P_marker[k1] != -i1-1) |
1086: { |
1087: P_diag_pass[pass][cnt_nz++] = k1; |
1088: P_marker[k1] = -i1-1; |
1089: } |
1090: } |
1091: j_start = P_offd_start[j1]; |
1092: j_end = j_start+P_offd_i[j1+1]; |
1093: for (k=j_start; k < j_end; k++) |
1094: { |
1095: k1 = P_offd_pass[pass-1][k]; |
1096: if (P_marker_offd[k1] != -i1-1) |
1097: { |
1098: P_offd_pass[pass][cnt_nz_offd++] = k1; |
1099: P_marker_offd[k1] = -i1-1; |
1100: } |
1101: } |
1102: } |
1103: } |
1104: for (j=S_offd_i[i1]; j < S_offd_i[i1+1]; j++) |
1105: { |
1106: j1 = S_offd_j[j]; |
1107: if (assigned_offd[j1] == pass-1) |
1108: { |
1109: j_start = Pext_start[j1]; |
1110: j_end = j_start+Pext_i[j1+1]; |
1111: for (k=j_start; k < j_end; k++) |
1112: { |
1113: k1 = Pext_pass[pass][k]; |
1114: if (k1 < 0) |
1115: { |
1116: if (P_marker[-k1-1] != -i1-1) |
1117: { |
1118: P_diag_pass[pass][cnt_nz++] = -k1-1; |
1119: P_marker[-k1-1] = -i1-1; |
1120: } |
1121: } |
1122: else if (P_marker_offd[k1] != -i1-1) |
1123: { |
1124: P_offd_pass[pass][cnt_nz_offd++] = k1; |
1125: P_marker_offd[k1] = -i1-1; |
Coverage (%) | Name | Source Location | Module |
---|---|---|---|
○95.55 | gomp_thread_start | team.c:130 | libgomp.so.1.0.0 |
○4.45 | GOMP_parallel | libgomp.h:985 | libgomp.so.1.0.0 |
Path / |
Metric | Value |
---|---|
CQA speedup if no scalar integer | 1.00 |
CQA speedup if FP arith vectorized | 1.00 |
CQA speedup if fully vectorized | 4.87 |
CQA speedup if no inter-iteration dependency | NA |
CQA speedup if next bottleneck killed | 1.30 |
Bottlenecks | P5, P6, P7, |
Function | hypre_BoomerAMGBuildMultipass._omp_fn.5 |
Source | par_multi_interp.c:1072-1075,par_multi_interp.c:1104-1104,par_multi_interp.c:1122-1122 |
Source loop unroll info | NA |
Source loop unroll confidence level | NA |
Unroll/vectorization loop type | NA |
Unroll factor | NA |
CQA cycles | 9.33 |
CQA cycles if no scalar integer | 9.33 |
CQA cycles if FP arith vectorized | 9.33 |
CQA cycles if fully vectorized | 1.92 |
Front-end cycles | 7.17 |
DIV/SQRT cycles | 2.75 |
P0 cycles | 2.75 |
P1 cycles | 2.50 |
P2 cycles | 2.50 |
P3 cycles | 2.50 |
P4 cycles | 9.33 |
P5 cycles | 9.33 |
P6 cycles | 9.33 |
P7 cycles | 0.00 |
P8 cycles | 0.00 |
P9 cycles | 0.00 |
P10 cycles | 0.00 |
P11 cycles | 0.00 |
P12 cycles | 0.00 |
P13 cycles | 0.00 |
Inter-iter dependencies cycles | NA |
FE+BE cycles (UFS) | NA |
Stall cycles (UFS) | NA |
Nb insns | 45.00 |
Nb uops | 43.00 |
Nb loads | 19.00 |
Nb stores | 10.00 |
Nb stack references | 12.00 |
FLOP/cycle | 0.00 |
Nb FLOP add-sub | 0.00 |
Nb FLOP mul | 0.00 |
Nb FLOP fma | 0.00 |
Nb FLOP div | 0.00 |
Nb FLOP rcp | 0.00 |
Nb FLOP sqrt | 0.00 |
Nb FLOP rsqrt | 0.00 |
Bytes/cycle | 24.86 |
Bytes prefetched | 0.00 |
Bytes loaded | 152.00 |
Bytes stored | 80.00 |
Stride 0 | NA |
Stride 1 | NA |
Stride n | NA |
Stride unknown | NA |
Stride indirect | NA |
Vectorization ratio all | 0.00 |
Vectorization ratio load | 0.00 |
Vectorization ratio store | 0.00 |
Vectorization ratio mul | NA |
Vectorization ratio add_sub | NA |
Vectorization ratio fma | NA |
Vectorization ratio div_sqrt | NA |
Vectorization ratio other | 0.00 |
Vector-efficiency ratio all | 12.50 |
Vector-efficiency ratio load | 12.50 |
Vector-efficiency ratio store | 12.50 |
Vector-efficiency ratio mul | NA |
Vector-efficiency ratio add_sub | NA |
Vector-efficiency ratio fma | NA |
Vector-efficiency ratio div_sqrt | NA |
Vector-efficiency ratio other | 12.50 |
Metric | Value |
---|---|
CQA speedup if no scalar integer | 1.00 |
CQA speedup if FP arith vectorized | 1.00 |
CQA speedup if fully vectorized | 4.87 |
CQA speedup if no inter-iteration dependency | NA |
CQA speedup if next bottleneck killed | 1.30 |
Bottlenecks | P5, P6, P7, |
Function | hypre_BoomerAMGBuildMultipass._omp_fn.5 |
Source | par_multi_interp.c:1072-1075,par_multi_interp.c:1104-1104,par_multi_interp.c:1122-1122 |
Source loop unroll info | NA |
Source loop unroll confidence level | NA |
Unroll/vectorization loop type | NA |
Unroll factor | NA |
CQA cycles | 9.33 |
CQA cycles if no scalar integer | 9.33 |
CQA cycles if FP arith vectorized | 9.33 |
CQA cycles if fully vectorized | 1.92 |
Front-end cycles | 7.17 |
DIV/SQRT cycles | 2.75 |
P0 cycles | 2.75 |
P1 cycles | 2.50 |
P2 cycles | 2.50 |
P3 cycles | 2.50 |
P4 cycles | 9.33 |
P5 cycles | 9.33 |
P6 cycles | 9.33 |
P7 cycles | 0.00 |
P8 cycles | 0.00 |
P9 cycles | 0.00 |
P10 cycles | 0.00 |
P11 cycles | 0.00 |
P12 cycles | 0.00 |
P13 cycles | 0.00 |
Inter-iter dependencies cycles | NA |
FE+BE cycles (UFS) | NA |
Stall cycles (UFS) | NA |
Nb insns | 45.00 |
Nb uops | 43.00 |
Nb loads | 19.00 |
Nb stores | 10.00 |
Nb stack references | 12.00 |
FLOP/cycle | 0.00 |
Nb FLOP add-sub | 0.00 |
Nb FLOP mul | 0.00 |
Nb FLOP fma | 0.00 |
Nb FLOP div | 0.00 |
Nb FLOP rcp | 0.00 |
Nb FLOP sqrt | 0.00 |
Nb FLOP rsqrt | 0.00 |
Bytes/cycle | 24.86 |
Bytes prefetched | 0.00 |
Bytes loaded | 152.00 |
Bytes stored | 80.00 |
Stride 0 | NA |
Stride 1 | NA |
Stride n | NA |
Stride unknown | NA |
Stride indirect | NA |
Vectorization ratio all | 0.00 |
Vectorization ratio load | 0.00 |
Vectorization ratio store | 0.00 |
Vectorization ratio mul | NA |
Vectorization ratio add_sub | NA |
Vectorization ratio fma | NA |
Vectorization ratio div_sqrt | NA |
Vectorization ratio other | 0.00 |
Vector-efficiency ratio all | 12.50 |
Vector-efficiency ratio load | 12.50 |
Vector-efficiency ratio store | 12.50 |
Vector-efficiency ratio mul | NA |
Vector-efficiency ratio add_sub | NA |
Vector-efficiency ratio fma | NA |
Vector-efficiency ratio div_sqrt | NA |
Vector-efficiency ratio other | 12.50 |
Path / |
Function | hypre_BoomerAMGBuildMultipass._omp_fn.5 |
Source file and lines | par_multi_interp.c:1072-1125 |
Module | libparcsr_ls.so |
nb instructions | 45 |
nb uops | 43 |
loop length | 230 |
used x86 registers | 11 |
used mmx registers | 0 |
used xmm registers | 0 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 12 |
micro-operation queue | 7.17 cycles |
front end | 7.17 cycles |
ALU0/BRU0 | ALU1 | ALU2 | ALU3 | BRU1 | AGU0 | AGU1 | AGU2 | FP0 | FP1 | FP2 | FP3 | FP4 | FP5 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 2.75 | 2.75 | 2.50 | 2.50 | 2.50 | 9.33 | 9.33 | 9.33 | 0.00 | 0.00 | 0.00 | 0.00 | 0.00 | 0.00 |
cycles | 2.75 | 2.75 | 2.50 | 2.50 | 2.50 | 9.33 | 9.33 | 9.33 | 0.00 | 0.00 | 0.00 | 0.00 | 0.00 | 0.00 |
Cycles executing div or sqrt instructions | NA |
Front-end | 7.17 |
Dispatch | 9.33 |
Overall L1 | 9.33 |
all | 0% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 0% |
all | 12% |
load | 12% |
store | 12% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 12% |
Instruction | Nb FU | ALU0/BRU0 | ALU1 | ALU2 | ALU3 | BRU1 | AGU0 | AGU1 | AGU2 | FP0 | FP1 | FP2 | FP3 | FP4 | FP5 | Latency | Recip. throughput | Vectorization |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
MOV -0xb0(%RBP),%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV -0x70(%RBP),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV -0xd0(%RBP),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV (%R13,%RCX,8),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
LEA 0x8(,%RDI,8),%RDX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
MOV (%R8,%RDI,8),%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %RDI,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 | N/A |
ADD %RDX,%R8 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
NOT %RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
MOV %R8,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV (%R8),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %R11,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
CMP %R8,%R11 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | scal (12.5%) |
JGE 7ffb0 <hypre_BoomerAMGBuildMultipass._omp_fn.5+0x1120> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 | N/A |
MOV %RDI,-0xe0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV %RBX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 | N/A |
MOV -0x88(%RBP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %RDX,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
JMP 7fcc1 <hypre_BoomerAMGBuildMultipass._omp_fn.5+0xe31> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | N/A |
NOPW (%RAX,%RAX,1) | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 | N/A |
MOV %RDI,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV -0xc0(%RBP),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | scal (12.5%) |
MOV %R13,%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 | scal (12.5%) |
MOV -0xe0(%RBP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV -0xc8(%RBP),%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV (%R11,%RDI,8),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
ADD %RDX,%R11 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
NOT %RDI | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
MOV (%R11),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %R11,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
CMP %RAX,%RCX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | scal (12.5%) |
JLE 8028a <hypre_BoomerAMGBuildMultipass._omp_fn.5+0x13fa> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 | N/A |
MOV %R9,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV -0x88(%RBP),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %R10,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV -0xd8(%RBP),%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
JMP 7fffc <hypre_BoomerAMGBuildMultipass._omp_fn.5+0x116c> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | N/A |
NOPL (%RAX,%RAX,1) | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 | N/A |
MOV %R8,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV -0x68(%RBP),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | scal (12.5%) |
MOV -0x90(%RBP),%R10 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | scal (12.5%) |
INCQ -0x70(%RBP) | 2 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 | scal (12.5%) |
MOV -0x70(%RBP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
CMP %RDI,-0xb8(%RBP) | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 | N/A |
JG 7fc50 <hypre_BoomerAMGBuildMultipass._omp_fn.5+0xdc0> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 | N/A |
Function | hypre_BoomerAMGBuildMultipass._omp_fn.5 |
Source file and lines | par_multi_interp.c:1072-1125 |
Module | libparcsr_ls.so |
nb instructions | 45 |
nb uops | 43 |
loop length | 230 |
used x86 registers | 11 |
used mmx registers | 0 |
used xmm registers | 0 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 12 |
micro-operation queue | 7.17 cycles |
front end | 7.17 cycles |
ALU0/BRU0 | ALU1 | ALU2 | ALU3 | BRU1 | AGU0 | AGU1 | AGU2 | FP0 | FP1 | FP2 | FP3 | FP4 | FP5 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 2.75 | 2.75 | 2.50 | 2.50 | 2.50 | 9.33 | 9.33 | 9.33 | 0.00 | 0.00 | 0.00 | 0.00 | 0.00 | 0.00 |
cycles | 2.75 | 2.75 | 2.50 | 2.50 | 2.50 | 9.33 | 9.33 | 9.33 | 0.00 | 0.00 | 0.00 | 0.00 | 0.00 | 0.00 |
Cycles executing div or sqrt instructions | NA |
Front-end | 7.17 |
Dispatch | 9.33 |
Overall L1 | 9.33 |
all | 0% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 0% |
all | 12% |
load | 12% |
store | 12% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | NA (no add-sub vectorizable/vectorized instructions) |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | NA (no div/sqrt vectorizable/vectorized instructions) |
other | 12% |
Instruction | Nb FU | ALU0/BRU0 | ALU1 | ALU2 | ALU3 | BRU1 | AGU0 | AGU1 | AGU2 | FP0 | FP1 | FP2 | FP3 | FP4 | FP5 | Latency | Recip. throughput | Vectorization |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
MOV -0xb0(%RBP),%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV -0x70(%RBP),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV -0xd0(%RBP),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV (%R13,%RCX,8),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
LEA 0x8(,%RDI,8),%RDX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
MOV (%R8,%RDI,8),%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %RDI,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 | N/A |
ADD %RDX,%R8 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
NOT %RAX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
MOV %R8,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV (%R8),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %R11,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
CMP %R8,%R11 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | scal (12.5%) |
JGE 7ffb0 <hypre_BoomerAMGBuildMultipass._omp_fn.5+0x1120> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 | N/A |
MOV %RDI,-0xe0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV %RBX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 | N/A |
MOV -0x88(%RBP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %RDX,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
JMP 7fcc1 <hypre_BoomerAMGBuildMultipass._omp_fn.5+0xe31> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | N/A |
NOPW (%RAX,%RAX,1) | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 | N/A |
MOV %RDI,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV -0xc0(%RBP),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | scal (12.5%) |
MOV %R13,%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 | scal (12.5%) |
MOV -0xe0(%RBP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV -0xc8(%RBP),%R11 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV (%R11,%RDI,8),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
ADD %RDX,%R11 | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
NOT %RDI | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | N/A |
MOV (%R11),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %R11,-0xc0(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
CMP %RAX,%RCX | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.25 | scal (12.5%) |
JLE 8028a <hypre_BoomerAMGBuildMultipass._omp_fn.5+0x13fa> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 | N/A |
MOV %R9,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV -0x88(%RBP),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
MOV %R10,-0x90(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV -0xd8(%RBP),%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
JMP 7fffc <hypre_BoomerAMGBuildMultipass._omp_fn.5+0x116c> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | N/A |
NOPL (%RAX,%RAX,1) | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.09 | N/A |
MOV %R8,-0x88(%RBP) | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 4 | 0.50 | scal (12.5%) |
MOV -0x68(%RBP),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | scal (12.5%) |
MOV -0x90(%RBP),%R10 | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | scal (12.5%) |
INCQ -0x70(%RBP) | 2 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 | scal (12.5%) |
MOV -0x70(%RBP),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 0.33 | N/A |
CMP %RDI,-0xb8(%RBP) | 1 | 0.25 | 0.25 | 0.25 | 0.25 | 0 | 0.33 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.33 | N/A |
JG 7fc50 <hypre_BoomerAMGBuildMultipass._omp_fn.5+0xdc0> | 1 | 0.50 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50-1 | N/A |