Function: hypre_CSRMatrixTranspose.extracted | Module: exec | Source: csr_matop.c:380-560 [...] | Coverage: 0.17% |
---|
Function: hypre_CSRMatrixTranspose.extracted | Module: exec | Source: csr_matop.c:380-560 [...] | Coverage: 0.17% |
---|
/scratch_na/users/xoserete/qaas_runs/171-172-8218/intel/AMG/build/AMG/AMG/seq_mv/csr_matop.c: 380 - 560 |
-------------------------------------------------------------------------------- |
380: return idx%dim1*dim2 + idx/dim1; |
[...] |
463: #pragma omp parallel |
464: #endif |
465: { |
466: HYPRE_Int num_threads = hypre_NumActiveThreads(); |
467: HYPRE_Int my_thread_num = hypre_GetThreadNum(); |
468: |
469: HYPRE_Int iBegin = hypre_CSRMatrixGetLoadBalancedPartitionBegin(A); |
470: HYPRE_Int iEnd = hypre_CSRMatrixGetLoadBalancedPartitionEnd(A); |
471: hypre_assert(iBegin <= iEnd); |
472: hypre_assert(iBegin >= 0 && iBegin <= num_rowsA); |
473: hypre_assert(iEnd >= 0 && iEnd <= num_rowsA); |
474: |
475: HYPRE_Int i, j; |
476: memset(bucket + my_thread_num*num_colsA, 0, sizeof(HYPRE_Int)*num_colsA); |
[...] |
483: for (j = A_i[iBegin]; j < A_i[iEnd]; ++j) { |
484: HYPRE_Int idx = A_j[j]; |
485: bucket[my_thread_num*num_colsA + idx]++; |
[...] |
496: for (i = my_thread_num*num_colsA + 1; i < (my_thread_num + 1)*num_colsA; ++i) { |
497: HYPRE_Int transpose_i = transpose_idx(i, num_threads, num_colsA); |
498: HYPRE_Int transpose_i_minus_1 = transpose_idx(i - 1, num_threads, num_colsA); |
499: |
500: bucket[transpose_i] += bucket[transpose_i_minus_1]; |
501: } |
502: |
503: #ifdef HYPRE_USING_OPENMP |
504: #pragma omp barrier |
505: #pragma omp master |
506: #endif |
507: { |
508: for (i = 1; i < num_threads; ++i) { |
509: HYPRE_Int j0 = num_colsA*i - 1, j1 = num_colsA*(i + 1) - 1; |
510: HYPRE_Int transpose_j0 = transpose_idx(j0, num_threads, num_colsA); |
511: HYPRE_Int transpose_j1 = transpose_idx(j1, num_threads, num_colsA); |
512: |
513: bucket[transpose_j1] += bucket[transpose_j0]; |
[...] |
520: if (my_thread_num > 0) { |
521: HYPRE_Int transpose_i0 = transpose_idx(num_colsA*my_thread_num - 1, num_threads, num_colsA); |
522: HYPRE_Int offset = bucket[transpose_i0]; |
523: |
524: for (i = my_thread_num*num_colsA; i < (my_thread_num + 1)*num_colsA - 1; ++i) { |
525: HYPRE_Int transpose_i = transpose_idx(i, num_threads, num_colsA); |
526: |
527: bucket[transpose_i] += offset; |
[...] |
539: if (data) { |
540: for (i = iEnd - 1; i >= iBegin; --i) { |
541: for (j = A_i[i + 1] - 1; j >= A_i[i]; --j) { |
542: HYPRE_Int idx = A_j[j]; |
543: --bucket[my_thread_num*num_colsA + idx]; |
544: |
545: HYPRE_Int offset = bucket[my_thread_num*num_colsA + idx]; |
546: |
547: AT_data[offset] = A_data[j]; |
548: AT_j[offset] = i; |
549: } |
550: } |
551: } |
552: else { |
553: for (i = iEnd - 1; i >= iBegin; --i) { |
554: for (j = A_i[i + 1] - 1; j >= A_i[i]; --j) { |
555: HYPRE_Int idx = A_j[j]; |
556: --bucket[my_thread_num*num_colsA + idx]; |
557: |
558: HYPRE_Int offset = bucket[my_thread_num*num_colsA + idx]; |
559: |
560: AT_j[offset] = i; |
0x4e9270 PUSH %RBP |
0x4e9271 MOV %RSP,%RBP |
0x4e9274 PUSH %R15 |
0x4e9276 PUSH %R14 |
0x4e9278 PUSH %R13 |
0x4e927a PUSH %R12 |
0x4e927c PUSH %RBX |
0x4e927d SUB $0x58,%RSP |
0x4e9281 MOV %R9,%RBX |
0x4e9284 MOV %R8,%R15 |
0x4e9287 MOV %RCX,-0x70(%RBP) |
0x4e928b MOV %RDX,%R12 |
0x4e928e MOV %RDI,-0x30(%RBP) |
0x4e9292 CALL 4f9c80 <hypre_NumActiveThreads> |
0x4e9297 MOV %RAX,%R13 |
0x4e929a CALL 4f9c90 <hypre_GetThreadNum> |
0x4e929f MOV %RAX,-0x48(%RBP) |
0x4e92a3 MOV %R12,%RDI |
0x4e92a6 CALL 4ebc90 <hypre_CSRMatrixGetLoadBalancedPartitionBegin> |
0x4e92ab MOV %RAX,%R14 |
0x4e92ae MOV %R12,%RDI |
0x4e92b1 CALL 4ebd10 <hypre_CSRMatrixGetLoadBalancedPartitionEnd> |
0x4e92b6 MOV %RAX,%RCX |
0x4e92b9 CMP %R14,%RAX |
0x4e92bc MOV %R14,%RAX |
0x4e92bf MOV %R14,-0x40(%RBP) |
0x4e92c3 MOV %RCX,-0x38(%RBP) |
0x4e92c7 JGE 4e92ff |
0x4e92c9 MOV 0x265930(%RIP),%RDI |
0x4e92d0 MOV $0x52769c,%ESI |
0x4e92d5 MOV $0x529f11,%EDX |
0x4e92da XOR %EAX,%EAX |
0x4e92dc CALL 4f8100 <hypre_fprintf> |
0x4e92e1 MOV $0x529e8a,%EDI |
0x4e92e6 MOV $0x1d7,%ESI |
0x4e92eb MOV $0x1,%EDX |
0x4e92f0 XOR %ECX,%ECX |
0x4e92f2 CALL 4faac0 <hypre_error_handler> |
0x4e92f7 MOV -0x38(%RBP),%RCX |
0x4e92fb MOV -0x40(%RBP),%R14 |
0x4e92ff MOV 0x18(%RBP),%R12 |
0x4e9303 TEST %R14,%R14 |
0x4e9306 JS 4e930d |
0x4e9308 CMP %R12,%R14 |
0x4e930b JLE 4e933f |
0x4e930d MOV 0x2658ec(%RIP),%RDI |
0x4e9314 MOV $0x52769c,%ESI |
0x4e9319 MOV $0x529f20,%EDX |
0x4e931e XOR %EAX,%EAX |
0x4e9320 CALL 4f8100 <hypre_fprintf> |
0x4e9325 MOV $0x529e8a,%EDI |
0x4e932a MOV $0x1d8,%ESI |
0x4e932f MOV $0x1,%EDX |
0x4e9334 XOR %ECX,%ECX |
0x4e9336 CALL 4faac0 <hypre_error_handler> |
0x4e933b MOV -0x38(%RBP),%RCX |
0x4e933f MOV 0x38(%RBP),%R14 |
0x4e9343 MOV 0x20(%RBP),%RAX |
0x4e9347 TEST %RCX,%RCX |
0x4e934a JS 4e9351 |
0x4e934c CMP %R12,%RCX |
0x4e934f JLE 4e9383 |
0x4e9351 MOV 0x2658a8(%RIP),%RDI |
0x4e9358 MOV $0x52769c,%ESI |
0x4e935d MOV $0x529f43,%EDX |
0x4e9362 XOR %EAX,%EAX |
0x4e9364 CALL 4f8100 <hypre_fprintf> |
0x4e9369 MOV $0x529e8a,%EDI |
0x4e936e MOV $0x1d9,%ESI |
0x4e9373 MOV $0x1,%EDX |
0x4e9378 XOR %ECX,%ECX |
0x4e937a CALL 4faac0 <hypre_error_handler> |
0x4e937f MOV 0x20(%RBP),%RAX |
0x4e9383 MOV -0x48(%RBP),%R12 |
0x4e9387 IMUL %RAX,%R12 |
0x4e938b LEA (%R14,%R12,8),%RDI |
0x4e938f LEA (,%RAX,8),%RDX |
0x4e9397 XOR %ESI,%ESI |
0x4e9399 CALL 5011c0 <_intel_fast_memset> |
0x4e939e MOV -0x38(%RBP),%RSI |
0x4e93a2 MOV 0x10(%RBP),%RDX |
0x4e93a6 MOV -0x40(%RBP),%RAX |
0x4e93aa MOV (%RBX,%RAX,8),%RAX |
0x4e93ae CMP (%RBX,%RSI,8),%RAX |
0x4e93b2 JGE 4e93d4 |
0x4e93b4 NOPW %CS:(%RAX,%RAX,1) |
(3444) 0x4e93c0 MOV (%RDX,%RAX,8),%RCX |
(3444) 0x4e93c4 ADD %R12,%RCX |
(3444) 0x4e93c7 INCQ (%R14,%RCX,8) |
(3444) 0x4e93cb INC %RAX |
(3444) 0x4e93ce CMP (%RBX,%RSI,8),%RAX |
(3444) 0x4e93d2 JL 4e93c0 |
0x4e93d4 MOV -0x30(%RBP),%RAX |
0x4e93d8 MOV (%RAX),%ESI |
0x4e93da MOV $0x74dab0,%EDI |
0x4e93df CALL 410130 <__kmpc_barrier@plt> |
0x4e93e4 MOV -0x48(%RBP),%RAX |
0x4e93e8 LEA 0x1(%RAX),%RCX |
0x4e93ec MOV 0x20(%RBP),%R11 |
0x4e93f0 IMUL %R11,%RCX |
0x4e93f4 LEA 0x1(%R12),%RAX |
0x4e93f9 MOV %RCX,-0x58(%RBP) |
0x4e93fd CMP %RCX,%RAX |
0x4e9400 JGE 4e94b1 |
0x4e9406 LEA -0x1(%R11),%R8 |
0x4e940a CMP $0x4,%R8 |
0x4e940e JAE 4e961c |
0x4e9414 MOV %R8,%R9 |
0x4e9417 AND $-0x4,%R9 |
0x4e941b CMP %R8,%R9 |
0x4e941e JAE 4e94b1 |
0x4e9424 LEA (%R12,%R9,1),%RDI |
0x4e9428 NOT %R9 |
0x4e942b ADD %R11,%R9 |
0x4e942e JMP 4e9456 |
(3442) 0x4e9430 MOV %RDI,%RAX |
(3442) 0x4e9433 CQTO |
(3442) 0x4e9435 IDIV %R13 |
(3442) 0x4e9438 IMUL %R11,%RDX |
(3442) 0x4e943c ADD %RAX,%RDX |
(3442) 0x4e943f MOV (%R14,%RDX,8),%RAX |
(3442) 0x4e9443 IMUL %R11,%RSI |
(3442) 0x4e9447 ADD %R8,%RSI |
(3442) 0x4e944a ADD %RAX,(%R14,%RSI,8) |
(3442) 0x4e944e MOV %RCX,%RDI |
(3442) 0x4e9451 DEC %R9 |
(3442) 0x4e9454 JE 4e94b1 |
(3442) 0x4e9456 LEA 0x1(%RDI),%RCX |
(3442) 0x4e945a MOV %RCX,%RAX |
(3442) 0x4e945d OR %R13,%RAX |
(3442) 0x4e9460 SHR $0x20,%RAX |
(3442) 0x4e9464 JE 4e9490 |
(3442) 0x4e9466 MOV %RCX,%RAX |
(3442) 0x4e9469 CQTO |
(3442) 0x4e946b IDIV %R13 |
(3442) 0x4e946e MOV %RDX,%RSI |
(3442) 0x4e9471 MOV %RAX,%R8 |
(3442) 0x4e9474 MOV %RDI,%RAX |
(3442) 0x4e9477 OR %R13,%RAX |
(3442) 0x4e947a SHR $0x20,%RAX |
(3442) 0x4e947e JNE 4e9430 |
(3442) 0x4e9480 JMP 4e94a8 |
0x4e9482 NOPW %CS:(%RAX,%RAX,1) |
(3442) 0x4e9490 MOV %ECX,%EAX |
(3442) 0x4e9492 XOR %EDX,%EDX |
(3442) 0x4e9494 DIV %R13D |
(3442) 0x4e9497 MOV %EDX,%ESI |
(3442) 0x4e9499 MOV %EAX,%R8D |
(3442) 0x4e949c MOV %RDI,%RAX |
(3442) 0x4e949f OR %R13,%RAX |
(3442) 0x4e94a2 SHR $0x20,%RAX |
(3442) 0x4e94a6 JNE 4e9430 |
(3442) 0x4e94a8 MOV %EDI,%EAX |
(3442) 0x4e94aa XOR %EDX,%EDX |
(3442) 0x4e94ac DIV %R13D |
(3442) 0x4e94af JMP 4e9438 |
0x4e94b1 MOV %R15,-0x68(%RBP) |
0x4e94b5 MOV -0x30(%RBP),%R15 |
0x4e94b9 MOV (%R15),%ESI |
0x4e94bc MOV $0x74dad0,%EDI |
0x4e94c1 CALL 410130 <__kmpc_barrier@plt> |
0x4e94c6 MOV (%R15),%ESI |
0x4e94c9 MOV -0x68(%RBP),%R15 |
0x4e94cd MOV $0x74daf0,%EDI |
0x4e94d2 XOR %EDX,%EDX |
0x4e94d4 CALL 4102a0 <__kmpc_masked@plt> |
0x4e94d9 CMP $0x1,%EAX |
0x4e94dc JNE 4e95ce |
0x4e94e2 CMP $0x1,%R13 |
0x4e94e6 MOV -0x30(%RBP),%R11 |
0x4e94ea JLE 4e95c1 |
0x4e94f0 LEA -0x1(%R13),%RAX |
0x4e94f4 MOV %RAX,-0x50(%RBP) |
0x4e94f8 CMP $0x4,%RAX |
0x4e94fc JAE 4e9a59 |
0x4e9502 MOV -0x50(%RBP),%RAX |
0x4e9506 MOV %RAX,%RCX |
0x4e9509 AND $-0x4,%RCX |
0x4e950d CMP %RAX,%RCX |
0x4e9510 JAE 4e95c1 |
0x4e9516 LEA 0x1(%RCX),%R9 |
0x4e951a MOV 0x20(%RBP),%RAX |
0x4e951e MOV %RAX,%RSI |
0x4e9521 IMUL %R9,%RSI |
0x4e9525 DEC %RSI |
0x4e9528 ADD $0x2,%RCX |
0x4e952c IMUL %RAX,%RCX |
0x4e9530 DEC %RCX |
0x4e9533 JMP 4e9570 |
0x4e9535 NOPW %CS:(%RAX,%RAX,1) |
(3440) 0x4e9540 MOV %RCX,%RAX |
(3440) 0x4e9543 CQTO |
(3440) 0x4e9545 IDIV %R13 |
(3440) 0x4e9548 MOV 0x20(%RBP),%R10 |
(3440) 0x4e954c IMUL %R10,%RDI |
(3440) 0x4e9550 ADD %R8,%RDI |
(3440) 0x4e9553 MOV (%R14,%RDI,8),%RDI |
(3440) 0x4e9557 IMUL %R10,%RDX |
(3440) 0x4e955b ADD %RAX,%RDX |
(3440) 0x4e955e ADD %RDI,(%R14,%RDX,8) |
(3440) 0x4e9562 ADD %R10,%RSI |
(3440) 0x4e9565 ADD %R10,%RCX |
(3440) 0x4e9568 INC %R9 |
(3440) 0x4e956b CMP %R9,%R13 |
(3440) 0x4e956e JE 4e95c1 |
(3440) 0x4e9570 MOV %RSI,%RAX |
(3440) 0x4e9573 OR %R13,%RAX |
(3440) 0x4e9576 SHR $0x20,%RAX |
(3440) 0x4e957a JE 4e95a0 |
(3440) 0x4e957c MOV %RSI,%RAX |
(3440) 0x4e957f CQTO |
(3440) 0x4e9581 IDIV %R13 |
(3440) 0x4e9584 MOV %RDX,%RDI |
(3440) 0x4e9587 MOV %RAX,%R8 |
(3440) 0x4e958a MOV %RCX,%RAX |
(3440) 0x4e958d OR %R13,%RAX |
(3440) 0x4e9590 SHR $0x20,%RAX |
(3440) 0x4e9594 JNE 4e9540 |
(3440) 0x4e9596 JMP 4e95b8 |
0x4e9598 NOPL (%RAX,%RAX,1) |
(3440) 0x4e95a0 MOV %ESI,%EAX |
(3440) 0x4e95a2 XOR %EDX,%EDX |
(3440) 0x4e95a4 DIV %R13D |
(3440) 0x4e95a7 MOV %EDX,%EDI |
(3440) 0x4e95a9 MOV %EAX,%R8D |
(3440) 0x4e95ac MOV %RCX,%RAX |
(3440) 0x4e95af OR %R13,%RAX |
(3440) 0x4e95b2 SHR $0x20,%RAX |
(3440) 0x4e95b6 JNE 4e9540 |
(3440) 0x4e95b8 MOV %ECX,%EAX |
(3440) 0x4e95ba XOR %EDX,%EDX |
(3440) 0x4e95bc DIV %R13D |
(3440) 0x4e95bf JMP 4e9548 |
0x4e95c1 MOV (%R11),%ESI |
0x4e95c4 MOV $0x74db10,%EDI |
0x4e95c9 CALL 4100c0 <__kmpc_end_masked@plt> |
0x4e95ce MOV -0x30(%RBP),%RAX |
0x4e95d2 MOV (%RAX),%ESI |
0x4e95d4 MOV $0x74db30,%EDI |
0x4e95d9 CALL 410130 <__kmpc_barrier@plt> |
0x4e95de CMPQ $0,-0x48(%RBP) |
0x4e95e3 MOV 0x20(%RBP),%R9 |
0x4e95e7 JLE 4e97b0 |
0x4e95ed LEA -0x1(%R12),%RAX |
0x4e95f2 MOV %RAX,%RCX |
0x4e95f5 OR %R13,%RCX |
0x4e95f8 SHR $0x20,%RCX |
0x4e95fc JE 4e9734 |
0x4e9602 CQTO |
0x4e9604 IDIV %R13 |
0x4e9607 MOV -0x58(%RBP),%RCX |
0x4e960b DEC %RCX |
0x4e960e CMP %RCX,%R12 |
0x4e9611 JL 4e9745 |
0x4e9617 JMP 4e97b0 |
0x4e961c MOV %R8,%R9 |
0x4e961f SHR $0x2,%R9 |
0x4e9623 MOV %R12,%RCX |
0x4e9626 JMP 4e964c |
0x4e9628 NOPL (%RAX,%RAX,1) |
(3443) 0x4e9630 MOV %RCX,%RAX |
(3443) 0x4e9633 CQTO |
(3443) 0x4e9635 IDIV %R13 |
(3443) 0x4e9638 IMUL %R11,%RDX |
(3443) 0x4e963c ADD %RAX,%RDX |
(3443) 0x4e963f ADD %R10,(%R14,%RDX,8) |
(3443) 0x4e9643 DEC %R9 |
(3443) 0x4e9646 JE 4e9414 |
(3443) 0x4e964c LEA 0x1(%RCX),%RAX |
(3443) 0x4e9650 MOV %RAX,%RDX |
(3443) 0x4e9653 OR %R13,%RDX |
(3443) 0x4e9656 SHR $0x20,%RDX |
(3443) 0x4e965a JE 4e9680 |
(3443) 0x4e965c CQTO |
(3443) 0x4e965e IDIV %R13 |
(3443) 0x4e9661 MOV %RDX,%RSI |
(3443) 0x4e9664 MOV %RAX,%RDI |
(3443) 0x4e9667 MOV %RCX,%RAX |
(3443) 0x4e966a OR %R13,%RAX |
(3443) 0x4e966d SHR $0x20,%RAX |
(3443) 0x4e9671 JE 4e9695 |
(3443) 0x4e9673 MOV %RCX,%RAX |
(3443) 0x4e9676 CQTO |
(3443) 0x4e9678 IDIV %R13 |
(3443) 0x4e967b JMP 4e969c |
0x4e967d NOPL (%RAX) |
(3443) 0x4e9680 XOR %EDX,%EDX |
(3443) 0x4e9682 DIV %R13D |
(3443) 0x4e9685 MOV %EDX,%ESI |
(3443) 0x4e9687 MOV %EAX,%EDI |
(3443) 0x4e9689 MOV %RCX,%RAX |
(3443) 0x4e968c OR %R13,%RAX |
(3443) 0x4e968f SHR $0x20,%RAX |
(3443) 0x4e9693 JNE 4e9673 |
(3443) 0x4e9695 MOV %ECX,%EAX |
(3443) 0x4e9697 XOR %EDX,%EDX |
(3443) 0x4e9699 DIV %R13D |
(3443) 0x4e969c IMUL %R11,%RDX |
(3443) 0x4e96a0 ADD %RAX,%RDX |
(3443) 0x4e96a3 MOV (%R14,%RDX,8),%R10 |
(3443) 0x4e96a7 IMUL %R11,%RSI |
(3443) 0x4e96ab ADD %RDI,%RSI |
(3443) 0x4e96ae ADD (%R14,%RSI,8),%R10 |
(3443) 0x4e96b2 MOV %R10,(%R14,%RSI,8) |
(3443) 0x4e96b6 LEA 0x2(%RCX),%RAX |
(3443) 0x4e96ba MOV %RAX,%RDX |
(3443) 0x4e96bd OR %R13,%RDX |
(3443) 0x4e96c0 SHR $0x20,%RDX |
(3443) 0x4e96c4 JE 4e96d0 |
(3443) 0x4e96c6 CQTO |
(3443) 0x4e96c8 IDIV %R13 |
(3443) 0x4e96cb JMP 4e96d5 |
0x4e96cd NOPL (%RAX) |
(3443) 0x4e96d0 XOR %EDX,%EDX |
(3443) 0x4e96d2 DIV %R13D |
(3443) 0x4e96d5 IMUL %R11,%RDX |
(3443) 0x4e96d9 ADD %RAX,%RDX |
(3443) 0x4e96dc ADD (%R14,%RDX,8),%R10 |
(3443) 0x4e96e0 MOV %R10,(%R14,%RDX,8) |
(3443) 0x4e96e4 LEA 0x3(%RCX),%RAX |
(3443) 0x4e96e8 MOV %RAX,%RDX |
(3443) 0x4e96eb OR %R13,%RDX |
(3443) 0x4e96ee SHR $0x20,%RDX |
(3443) 0x4e96f2 JE 4e9700 |
(3443) 0x4e96f4 CQTO |
(3443) 0x4e96f6 IDIV %R13 |
(3443) 0x4e96f9 JMP 4e9705 |
0x4e96fb NOPL (%RAX,%RAX,1) |
(3443) 0x4e9700 XOR %EDX,%EDX |
(3443) 0x4e9702 DIV %R13D |
(3443) 0x4e9705 IMUL %R11,%RDX |
(3443) 0x4e9709 ADD %RAX,%RDX |
(3443) 0x4e970c ADD (%R14,%RDX,8),%R10 |
(3443) 0x4e9710 MOV %R10,(%R14,%RDX,8) |
(3443) 0x4e9714 ADD $0x4,%RCX |
(3443) 0x4e9718 MOV %RCX,%RAX |
(3443) 0x4e971b OR %R13,%RAX |
(3443) 0x4e971e SHR $0x20,%RAX |
(3443) 0x4e9722 JNE 4e9630 |
(3443) 0x4e9728 MOV %ECX,%EAX |
(3443) 0x4e972a XOR %EDX,%EDX |
(3443) 0x4e972c DIV %R13D |
(3443) 0x4e972f JMP 4e9638 |
0x4e9734 XOR %EDX,%EDX |
0x4e9736 DIV %R13D |
0x4e9739 MOV -0x58(%RBP),%RCX |
0x4e973d DEC %RCX |
0x4e9740 CMP %RCX,%R12 |
0x4e9743 JGE 4e97b0 |
0x4e9745 IMUL %R9,%RDX |
0x4e9749 ADD %RAX,%RDX |
0x4e974c MOV (%R14,%RDX,8),%RSI |
0x4e9750 LEA -0x1(%R9),%RDI |
0x4e9754 CMP $0x8,%RDI |
0x4e9758 JAE 4e9847 |
0x4e975e MOV %RDI,%R8 |
0x4e9761 AND $-0x8,%R8 |
0x4e9765 CMP %RDI,%R8 |
0x4e9768 JAE 4e97b0 |
0x4e976a LEA (%R12,%R8,1),%RCX |
0x4e976e NOT %R8 |
0x4e9771 ADD %R9,%R8 |
0x4e9774 JMP 4e979b |
0x4e9776 NOPW %CS:(%RAX,%RAX,1) |
(3438) 0x4e9780 MOV %RCX,%RAX |
(3438) 0x4e9783 CQTO |
(3438) 0x4e9785 IDIV %R13 |
(3438) 0x4e9788 IMUL %R9,%RDX |
(3438) 0x4e978c ADD %RAX,%RDX |
(3438) 0x4e978f ADD %RSI,(%R14,%RDX,8) |
(3438) 0x4e9793 INC %RCX |
(3438) 0x4e9796 DEC %R8 |
(3438) 0x4e9799 JE 4e97b0 |
(3438) 0x4e979b MOV %RCX,%RAX |
(3438) 0x4e979e OR %R13,%RAX |
(3438) 0x4e97a1 SHR $0x20,%RAX |
(3438) 0x4e97a5 JNE 4e9780 |
(3438) 0x4e97a7 MOV %ECX,%EAX |
(3438) 0x4e97a9 XOR %EDX,%EDX |
(3438) 0x4e97ab DIV %R13D |
(3438) 0x4e97ae JMP 4e9788 |
0x4e97b0 MOV 0x30(%RBP),%R13 |
0x4e97b4 MOV -0x30(%RBP),%RAX |
0x4e97b8 MOV (%RAX),%ESI |
0x4e97ba MOV $0x74db50,%EDI |
0x4e97bf CALL 410130 <__kmpc_barrier@plt> |
0x4e97c4 CMPQ $0,-0x70(%RBP) |
0x4e97c9 JE 4e980b |
0x4e97cb MOV -0x40(%RBP),%R10 |
0x4e97cf MOV -0x38(%RBP),%RSI |
0x4e97d3 CMP %R10,%RSI |
0x4e97d6 MOV 0x10(%RBP),%R11 |
0x4e97da JLE 4e9cd6 |
0x4e97e0 MOV 0x28(%RBP),%RAX |
0x4e97e4 MOV (%RBX,%RSI,8),%RDX |
0x4e97e8 MOV %ESI,%ECX |
0x4e97ea SUB %R10D,%ECX |
0x4e97ed LEA 0x1(%R10),%R8 |
0x4e97f1 TEST $0x1,%CL |
0x4e97f4 JNE 4e99ec |
0x4e97fa MOV %RSI,%RCX |
0x4e97fd CMP %R8,%RSI |
0x4e9800 JNE 4e9cf5 |
0x4e9806 JMP 4e9cd6 |
0x4e980b MOV -0x40(%RBP),%R9 |
0x4e980f MOV -0x38(%RBP),%R11 |
0x4e9813 CMP %R9,%R11 |
0x4e9816 MOV 0x10(%RBP),%R10 |
0x4e981a JLE 4e9cd6 |
0x4e9820 MOV (%RBX,%R11,8),%RCX |
0x4e9824 MOV %R11D,%EAX |
0x4e9827 SUB %R9D,%EAX |
0x4e982a LEA 0x1(%R9),%RDX |
0x4e982e TEST $0x1,%AL |
0x4e9830 JNE 4e9be3 |
0x4e9836 MOV %R11,%RAX |
0x4e9839 CMP %RDX,%R11 |
0x4e983c JNE 4e9c49 |
0x4e9842 JMP 4e9cd6 |
0x4e9847 MOV %RDI,%R8 |
0x4e984a SHR $0x3,%R8 |
0x4e984e LEA 0x7(%R12),%RCX |
0x4e9853 JMP 4e9880 |
0x4e9855 NOPW %CS:(%RAX,%RAX,1) |
(3439) 0x4e9860 MOV %RCX,%RAX |
(3439) 0x4e9863 CQTO |
(3439) 0x4e9865 IDIV %R13 |
(3439) 0x4e9868 IMUL %R9,%RDX |
(3439) 0x4e986c ADD %RAX,%RDX |
(3439) 0x4e986f ADD %RSI,(%R14,%RDX,8) |
(3439) 0x4e9873 ADD $0x8,%RCX |
(3439) 0x4e9877 DEC %R8 |
(3439) 0x4e987a JE 4e975e |
(3439) 0x4e9880 LEA -0x7(%RCX),%RAX |
(3439) 0x4e9884 MOV %RAX,%RDX |
(3439) 0x4e9887 OR %R13,%RDX |
(3439) 0x4e988a SHR $0x20,%RDX |
(3439) 0x4e988e JE 4e98a0 |
(3439) 0x4e9890 CQTO |
(3439) 0x4e9892 IDIV %R13 |
(3439) 0x4e9895 JMP 4e98a5 |
0x4e9897 NOPW (%RAX,%RAX,1) |
(3439) 0x4e98a0 XOR %EDX,%EDX |
(3439) 0x4e98a2 DIV %R13D |
(3439) 0x4e98a5 IMUL %R9,%RDX |
(3439) 0x4e98a9 ADD %RAX,%RDX |
(3439) 0x4e98ac ADD %RSI,(%R14,%RDX,8) |
(3439) 0x4e98b0 LEA -0x6(%RCX),%RAX |
(3439) 0x4e98b4 MOV %RAX,%RDX |
(3439) 0x4e98b7 OR %R13,%RDX |
(3439) 0x4e98ba SHR $0x20,%RDX |
(3439) 0x4e98be JE 4e98d0 |
(3439) 0x4e98c0 CQTO |
(3439) 0x4e98c2 IDIV %R13 |
(3439) 0x4e98c5 JMP 4e98d5 |
0x4e98c7 NOPW (%RAX,%RAX,1) |
(3439) 0x4e98d0 XOR %EDX,%EDX |
(3439) 0x4e98d2 DIV %R13D |
(3439) 0x4e98d5 IMUL %R9,%RDX |
(3439) 0x4e98d9 ADD %RAX,%RDX |
(3439) 0x4e98dc ADD %RSI,(%R14,%RDX,8) |
(3439) 0x4e98e0 LEA -0x5(%RCX),%RAX |
(3439) 0x4e98e4 MOV %RAX,%RDX |
(3439) 0x4e98e7 OR %R13,%RDX |
(3439) 0x4e98ea SHR $0x20,%RDX |
(3439) 0x4e98ee JE 4e9900 |
(3439) 0x4e98f0 CQTO |
(3439) 0x4e98f2 IDIV %R13 |
(3439) 0x4e98f5 JMP 4e9905 |
0x4e98f7 NOPW (%RAX,%RAX,1) |
(3439) 0x4e9900 XOR %EDX,%EDX |
(3439) 0x4e9902 DIV %R13D |
(3439) 0x4e9905 IMUL %R9,%RDX |
(3439) 0x4e9909 ADD %RAX,%RDX |
(3439) 0x4e990c ADD %RSI,(%R14,%RDX,8) |
(3439) 0x4e9910 LEA -0x4(%RCX),%RAX |
(3439) 0x4e9914 MOV %RAX,%RDX |
(3439) 0x4e9917 OR %R13,%RDX |
(3439) 0x4e991a SHR $0x20,%RDX |
(3439) 0x4e991e JE 4e9930 |
(3439) 0x4e9920 CQTO |
(3439) 0x4e9922 IDIV %R13 |
(3439) 0x4e9925 JMP 4e9935 |
0x4e9927 NOPW (%RAX,%RAX,1) |
(3439) 0x4e9930 XOR %EDX,%EDX |
(3439) 0x4e9932 DIV %R13D |
(3439) 0x4e9935 IMUL %R9,%RDX |
(3439) 0x4e9939 ADD %RAX,%RDX |
(3439) 0x4e993c ADD %RSI,(%R14,%RDX,8) |
(3439) 0x4e9940 LEA -0x3(%RCX),%RAX |
(3439) 0x4e9944 MOV %RAX,%RDX |
(3439) 0x4e9947 OR %R13,%RDX |
(3439) 0x4e994a SHR $0x20,%RDX |
(3439) 0x4e994e JE 4e9960 |
(3439) 0x4e9950 CQTO |
(3439) 0x4e9952 IDIV %R13 |
(3439) 0x4e9955 JMP 4e9965 |
0x4e9957 NOPW (%RAX,%RAX,1) |
(3439) 0x4e9960 XOR %EDX,%EDX |
(3439) 0x4e9962 DIV %R13D |
(3439) 0x4e9965 IMUL %R9,%RDX |
(3439) 0x4e9969 ADD %RAX,%RDX |
(3439) 0x4e996c ADD %RSI,(%R14,%RDX,8) |
(3439) 0x4e9970 LEA -0x2(%RCX),%RAX |
(3439) 0x4e9974 MOV %RAX,%RDX |
(3439) 0x4e9977 OR %R13,%RDX |
(3439) 0x4e997a SHR $0x20,%RDX |
(3439) 0x4e997e JE 4e9990 |
(3439) 0x4e9980 CQTO |
(3439) 0x4e9982 IDIV %R13 |
(3439) 0x4e9985 JMP 4e9995 |
0x4e9987 NOPW (%RAX,%RAX,1) |
(3439) 0x4e9990 XOR %EDX,%EDX |
(3439) 0x4e9992 DIV %R13D |
(3439) 0x4e9995 IMUL %R9,%RDX |
(3439) 0x4e9999 ADD %RAX,%RDX |
(3439) 0x4e999c ADD %RSI,(%R14,%RDX,8) |
(3439) 0x4e99a0 LEA -0x1(%RCX),%RAX |
(3439) 0x4e99a4 MOV %RAX,%RDX |
(3439) 0x4e99a7 OR %R13,%RDX |
(3439) 0x4e99aa SHR $0x20,%RDX |
(3439) 0x4e99ae JE 4e99c0 |
(3439) 0x4e99b0 CQTO |
(3439) 0x4e99b2 IDIV %R13 |
(3439) 0x4e99b5 JMP 4e99c5 |
0x4e99b7 NOPW (%RAX,%RAX,1) |
(3439) 0x4e99c0 XOR %EDX,%EDX |
(3439) 0x4e99c2 DIV %R13D |
(3439) 0x4e99c5 IMUL %R9,%RDX |
(3439) 0x4e99c9 ADD %RAX,%RDX |
(3439) 0x4e99cc ADD %RSI,(%R14,%RDX,8) |
(3439) 0x4e99d0 MOV %RCX,%RAX |
(3439) 0x4e99d3 OR %R13,%RAX |
(3439) 0x4e99d6 SHR $0x20,%RAX |
(3439) 0x4e99da JNE 4e9860 |
(3439) 0x4e99e0 MOV %ECX,%EAX |
(3439) 0x4e99e2 XOR %EDX,%EDX |
(3439) 0x4e99e4 DIV %R13D |
(3439) 0x4e99e7 JMP 4e9868 |
0x4e99ec LEA -0x1(%RSI),%RCX |
0x4e99f0 MOV -0x8(%RBX,%RSI,8),%RDI |
0x4e99f5 CMP %RDI,%RDX |
0x4e99f8 JLE 4e9cca |
0x4e99fe MOV %R8,-0x30(%RBP) |
0x4e9a02 MOV -0x38(%RBP),%RSI |
0x4e9a06 NOPW %CS:(%RAX,%RAX,1) |
(3437) 0x4e9a10 MOV -0x8(%R11,%RDX,8),%RDI |
(3437) 0x4e9a15 ADD %R12,%RDI |
(3437) 0x4e9a18 MOV (%R14,%RDI,8),%R8 |
(3437) 0x4e9a1c LEA -0x1(%R8),%R9 |
(3437) 0x4e9a20 MOV %R9,(%R14,%RDI,8) |
(3437) 0x4e9a24 VMOVSD -0x8(%R15,%RDX,8),%XMM0 |
(3437) 0x4e9a2b VMOVSD %XMM0,-0x8(%RAX,%R8,8) |
(3437) 0x4e9a32 DEC %RDX |
(3437) 0x4e9a35 MOV %RCX,-0x8(%R13,%R8,8) |
(3437) 0x4e9a3a MOV -0x8(%RBX,%RSI,8),%RDI |
(3437) 0x4e9a3f CMP %RDI,%RDX |
(3437) 0x4e9a42 JG 4e9a10 |
0x4e9a44 MOV %RDI,%RDX |
0x4e9a47 MOV -0x30(%RBP),%R8 |
0x4e9a4b CMP %R8,%RSI |
0x4e9a4e JNE 4e9cf5 |
0x4e9a54 JMP 4e9cd6 |
0x4e9a59 MOV -0x50(%RBP),%RDX |
0x4e9a5d SHR $0x2,%RDX |
0x4e9a61 MOV 0x20(%RBP),%RAX |
0x4e9a65 LEA (,%RAX,4),%RCX |
0x4e9a6d MOV %RCX,-0x78(%RBP) |
0x4e9a71 LEA (%RAX,%RAX,4),%RCX |
0x4e9a75 DEC %RCX |
0x4e9a78 LEA -0x1(,%RAX,4),%RSI |
0x4e9a80 LEA (%RAX,%RAX,2),%RDI |
0x4e9a84 DEC %RDI |
0x4e9a87 LEA -0x1(,%RAX,2),%R8 |
0x4e9a8f LEA -0x1(%RAX),%R9 |
0x4e9a93 JMP 4e9ad4 |
0x4e9a95 NOPW %CS:(%RAX,%RAX,1) |
(3441) 0x4e9aa0 MOV %RCX,%RAX |
(3441) 0x4e9aa3 CQTO |
(3441) 0x4e9aa5 IDIV %R13 |
(3441) 0x4e9aa8 IMUL 0x20(%RBP),%RDX |
(3441) 0x4e9aad ADD %RAX,%RDX |
(3441) 0x4e9ab0 ADD %R10,(%R14,%RDX,8) |
(3441) 0x4e9ab4 MOV -0x78(%RBP),%RAX |
(3441) 0x4e9ab8 ADD %RAX,%RCX |
(3441) 0x4e9abb ADD %RAX,%RSI |
(3441) 0x4e9abe ADD %RAX,%RDI |
(3441) 0x4e9ac1 ADD %RAX,%R8 |
(3441) 0x4e9ac4 ADD %RAX,%R9 |
(3441) 0x4e9ac7 MOV -0x80(%RBP),%RDX |
(3441) 0x4e9acb DEC %RDX |
(3441) 0x4e9ace JE 4e9502 |
(3441) 0x4e9ad4 MOV %R9,%RAX |
(3441) 0x4e9ad7 OR %R13,%RAX |
(3441) 0x4e9ada SHR $0x20,%RAX |
(3441) 0x4e9ade MOV %RDX,-0x80(%RBP) |
(3441) 0x4e9ae2 JE 4e9b00 |
(3441) 0x4e9ae4 MOV %R9,%RAX |
(3441) 0x4e9ae7 CQTO |
(3441) 0x4e9ae9 IDIV %R13 |
(3441) 0x4e9aec MOV %RDX,%R10 |
(3441) 0x4e9aef JMP 4e9b0b |
0x4e9af1 NOPW %CS:(%RAX,%RAX,1) |
(3441) 0x4e9b00 MOV %R9D,%EAX |
(3441) 0x4e9b03 XOR %EDX,%EDX |
(3441) 0x4e9b05 DIV %R13D |
(3441) 0x4e9b08 MOV %EDX,%R10D |
(3441) 0x4e9b0b MOV %RAX,-0x60(%RBP) |
(3441) 0x4e9b0f MOV %R8,%RAX |
(3441) 0x4e9b12 OR %R13,%RAX |
(3441) 0x4e9b15 SHR $0x20,%RAX |
(3441) 0x4e9b19 JE 4e9b30 |
(3441) 0x4e9b1b MOV %R8,%RAX |
(3441) 0x4e9b1e CQTO |
(3441) 0x4e9b20 IDIV %R13 |
(3441) 0x4e9b23 JMP 4e9b38 |
0x4e9b25 NOPW %CS:(%RAX,%RAX,1) |
(3441) 0x4e9b30 MOV %R8D,%EAX |
(3441) 0x4e9b33 XOR %EDX,%EDX |
(3441) 0x4e9b35 DIV %R13D |
(3441) 0x4e9b38 MOV 0x20(%RBP),%R11 |
(3441) 0x4e9b3c IMUL %R11,%R10 |
(3441) 0x4e9b40 ADD -0x60(%RBP),%R10 |
(3441) 0x4e9b44 MOV (%R14,%R10,8),%R10 |
(3441) 0x4e9b48 IMUL %R11,%RDX |
(3441) 0x4e9b4c ADD %RAX,%RDX |
(3441) 0x4e9b4f ADD (%R14,%RDX,8),%R10 |
(3441) 0x4e9b53 MOV %R10,(%R14,%RDX,8) |
(3441) 0x4e9b57 MOV %RDI,%RAX |
(3441) 0x4e9b5a OR %R13,%RAX |
(3441) 0x4e9b5d SHR $0x20,%RAX |
(3441) 0x4e9b61 JE 4e9b70 |
(3441) 0x4e9b63 MOV %RDI,%RAX |
(3441) 0x4e9b66 CQTO |
(3441) 0x4e9b68 IDIV %R13 |
(3441) 0x4e9b6b JMP 4e9b77 |
0x4e9b6d NOPL (%RAX) |
(3441) 0x4e9b70 MOV %EDI,%EAX |
(3441) 0x4e9b72 XOR %EDX,%EDX |
(3441) 0x4e9b74 DIV %R13D |
(3441) 0x4e9b77 MOV -0x30(%RBP),%R11 |
(3441) 0x4e9b7b IMUL 0x20(%RBP),%RDX |
(3441) 0x4e9b80 ADD %RAX,%RDX |
(3441) 0x4e9b83 ADD (%R14,%RDX,8),%R10 |
(3441) 0x4e9b87 MOV %R10,(%R14,%RDX,8) |
(3441) 0x4e9b8b MOV %RSI,%RAX |
(3441) 0x4e9b8e OR %R13,%RAX |
(3441) 0x4e9b91 SHR $0x20,%RAX |
(3441) 0x4e9b95 JE 4e9bb0 |
(3441) 0x4e9b97 MOV %RSI,%RAX |
(3441) 0x4e9b9a CQTO |
(3441) 0x4e9b9c IDIV %R13 |
(3441) 0x4e9b9f JMP 4e9bb7 |
0x4e9ba1 NOPW %CS:(%RAX,%RAX,1) |
(3441) 0x4e9bb0 MOV %ESI,%EAX |
(3441) 0x4e9bb2 XOR %EDX,%EDX |
(3441) 0x4e9bb4 DIV %R13D |
(3441) 0x4e9bb7 IMUL 0x20(%RBP),%RDX |
(3441) 0x4e9bbc ADD %RAX,%RDX |
(3441) 0x4e9bbf ADD (%R14,%RDX,8),%R10 |
(3441) 0x4e9bc3 MOV %R10,(%R14,%RDX,8) |
(3441) 0x4e9bc7 MOV %RCX,%RAX |
(3441) 0x4e9bca OR %R13,%RAX |
(3441) 0x4e9bcd SHR $0x20,%RAX |
(3441) 0x4e9bd1 JNE 4e9aa0 |
(3441) 0x4e9bd7 MOV %ECX,%EAX |
(3441) 0x4e9bd9 XOR %EDX,%EDX |
(3441) 0x4e9bdb DIV %R13D |
(3441) 0x4e9bde JMP 4e9aa8 |
0x4e9be3 LEA -0x1(%R11),%RAX |
0x4e9be7 MOV -0x8(%RBX,%R11,8),%RSI |
0x4e9bec CMP %RSI,%RCX |
0x4e9bef JLE 4e9c26 |
0x4e9bf1 NOPW %CS:(%RAX,%RAX,1) |
(3433) 0x4e9c00 MOV -0x8(%R10,%RCX,8),%RSI |
(3433) 0x4e9c05 ADD %R12,%RSI |
(3433) 0x4e9c08 DEC %RCX |
(3433) 0x4e9c0b MOV (%R14,%RSI,8),%RDI |
(3433) 0x4e9c0f LEA -0x1(%RDI),%R8 |
(3433) 0x4e9c13 MOV %R8,(%R14,%RSI,8) |
(3433) 0x4e9c17 MOV %RAX,-0x8(%R13,%RDI,8) |
(3433) 0x4e9c1c MOV -0x8(%RBX,%R11,8),%RSI |
(3433) 0x4e9c21 CMP %RSI,%RCX |
(3433) 0x4e9c24 JG 4e9c00 |
0x4e9c26 MOV %RSI,%RCX |
0x4e9c29 CMP %RDX,%R11 |
0x4e9c2c JNE 4e9c49 |
0x4e9c2e JMP 4e9cd6 |
0x4e9c33 NOPW %CS:(%RAX,%RAX,1) |
(3430) 0x4e9c40 CMP %R9,%RAX |
(3430) 0x4e9c43 JLE 4e9cd6 |
(3430) 0x4e9c49 MOV -0x8(%RBX,%RAX,8),%RDX |
(3430) 0x4e9c4e CMP %RDX,%RCX |
(3430) 0x4e9c51 JLE 4e9c86 |
(3430) 0x4e9c53 LEA -0x1(%RAX),%RSI |
(3430) 0x4e9c57 NOPW (%RAX,%RAX,1) |
(3432) 0x4e9c60 MOV -0x8(%R10,%RCX,8),%RDX |
(3432) 0x4e9c65 ADD %R12,%RDX |
(3432) 0x4e9c68 DEC %RCX |
(3432) 0x4e9c6b MOV (%R14,%RDX,8),%RDI |
(3432) 0x4e9c6f LEA -0x1(%RDI),%R8 |
(3432) 0x4e9c73 MOV %R8,(%R14,%RDX,8) |
(3432) 0x4e9c77 MOV %RSI,-0x8(%R13,%RDI,8) |
(3432) 0x4e9c7c MOV -0x8(%RBX,%RAX,8),%RDX |
(3432) 0x4e9c81 CMP %RDX,%RCX |
(3432) 0x4e9c84 JG 4e9c60 |
(3430) 0x4e9c86 MOV -0x10(%RBX,%RAX,8),%RCX |
(3430) 0x4e9c8b ADD $-0x2,%RAX |
(3430) 0x4e9c8f CMP %RCX,%RDX |
(3430) 0x4e9c92 JLE 4e9c40 |
(3430) 0x4e9c94 NOPW %CS:(%RAX,%RAX,1) |
(3431) 0x4e9ca0 MOV -0x8(%R10,%RDX,8),%RCX |
(3431) 0x4e9ca5 ADD %R12,%RCX |
(3431) 0x4e9ca8 DEC %RDX |
(3431) 0x4e9cab MOV (%R14,%RCX,8),%RSI |
(3431) 0x4e9caf LEA -0x1(%RSI),%RDI |
(3431) 0x4e9cb3 MOV %RDI,(%R14,%RCX,8) |
(3431) 0x4e9cb7 MOV %RAX,-0x8(%R13,%RSI,8) |
(3431) 0x4e9cbc MOV (%RBX,%RAX,8),%RCX |
(3431) 0x4e9cc0 CMP %RCX,%RDX |
(3431) 0x4e9cc3 JG 4e9ca0 |
(3430) 0x4e9cc5 JMP 4e9c40 |
0x4e9cca MOV %RDI,%RDX |
0x4e9ccd MOV -0x38(%RBP),%RSI |
0x4e9cd1 CMP %R8,%RSI |
0x4e9cd4 JNE 4e9cf5 |
0x4e9cd6 ADD $0x58,%RSP |
0x4e9cda POP %RBX |
0x4e9cdb POP %R12 |
0x4e9cdd POP %R13 |
0x4e9cdf POP %R14 |
0x4e9ce1 POP %R15 |
0x4e9ce3 POP %RBP |
0x4e9ce4 RET |
0x4e9ce5 NOPW %CS:(%RAX,%RAX,1) |
(3434) 0x4e9cf0 CMP %R10,%RCX |
(3434) 0x4e9cf3 JLE 4e9cd6 |
(3434) 0x4e9cf5 MOV -0x8(%RBX,%RCX,8),%RSI |
(3434) 0x4e9cfa CMP %RSI,%RDX |
(3434) 0x4e9cfd JLE 4e9d44 |
(3434) 0x4e9cff LEA -0x1(%RCX),%RDI |
(3434) 0x4e9d03 NOPW %CS:(%RAX,%RAX,1) |
(3436) 0x4e9d10 MOV -0x8(%R11,%RDX,8),%RSI |
(3436) 0x4e9d15 ADD %R12,%RSI |
(3436) 0x4e9d18 MOV (%R14,%RSI,8),%R8 |
(3436) 0x4e9d1c LEA -0x1(%R8),%R9 |
(3436) 0x4e9d20 MOV %R9,(%R14,%RSI,8) |
(3436) 0x4e9d24 VMOVSD -0x8(%R15,%RDX,8),%XMM0 |
(3436) 0x4e9d2b VMOVSD %XMM0,-0x8(%RAX,%R8,8) |
(3436) 0x4e9d32 DEC %RDX |
(3436) 0x4e9d35 MOV %RDI,-0x8(%R13,%R8,8) |
(3436) 0x4e9d3a MOV -0x8(%RBX,%RCX,8),%RSI |
(3436) 0x4e9d3f CMP %RSI,%RDX |
(3436) 0x4e9d42 JG 4e9d10 |
(3434) 0x4e9d44 MOV -0x10(%RBX,%RCX,8),%RDX |
(3434) 0x4e9d49 ADD $-0x2,%RCX |
(3434) 0x4e9d4d CMP %RDX,%RSI |
(3434) 0x4e9d50 JLE 4e9cf0 |
(3434) 0x4e9d52 NOPW %CS:(%RAX,%RAX,1) |
(3435) 0x4e9d60 MOV -0x8(%R11,%RSI,8),%RDX |
(3435) 0x4e9d65 ADD %R12,%RDX |
(3435) 0x4e9d68 MOV (%R14,%RDX,8),%RDI |
(3435) 0x4e9d6c LEA -0x1(%RDI),%R8 |
(3435) 0x4e9d70 MOV %R8,(%R14,%RDX,8) |
(3435) 0x4e9d74 VMOVSD -0x8(%R15,%RSI,8),%XMM0 |
(3435) 0x4e9d7b VMOVSD %XMM0,-0x8(%RAX,%RDI,8) |
(3435) 0x4e9d81 DEC %RSI |
(3435) 0x4e9d84 MOV %RCX,-0x8(%R13,%RDI,8) |
(3435) 0x4e9d89 MOV (%RBX,%RCX,8),%RDX |
(3435) 0x4e9d8d CMP %RDX,%RSI |
(3435) 0x4e9d90 JG 4e9d60 |
(3434) 0x4e9d92 JMP 4e9cf0 |
0x4e9d97 NOPW (%RAX,%RAX,1) |
Path / |
Source file and lines | csr_matop.c:380-560 |
Module | exec |
nb instructions | 300 |
nb uops | 324 |
loop length | 1304 |
used x86 registers | 16 |
used mmx registers | 0 |
used xmm registers | 0 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 15 |
micro-operation queue | 54.17 cycles |
front end | 54.17 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 22.10 | 22.00 | 20.67 | 20.67 | 16.50 | 22.00 | 21.90 | 16.50 | 16.50 | 16.50 | 22.00 | 20.67 |
cycles | 22.10 | 23.40 | 20.67 | 20.67 | 16.50 | 22.00 | 21.90 | 16.50 | 16.50 | 16.50 | 22.00 | 20.67 |
Cycles executing div or sqrt instructions | 16.00 |
FE+BE cycles | 51.00-51.06 |
Stall cycles | 0.00 |
Front-end | 54.17 |
Dispatch | 23.40 |
DIV/SQRT | 16.00 |
Overall L1 | 54.17 |
all | 0% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 0% |
other | 0% |
all | 11% |
load | 12% |
store | 12% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 9% |
other | 10% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
SUB $0x58,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R9,%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R8,%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %RCX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDX,%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %RDI,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CALL 4f9c80 <hypre_NumActiveThreads> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 4f9c90 <hypre_GetThreadNum> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R12,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 4ebc90 <hypre_CSRMatrixGetLoadBalancedPartitionBegin> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R12,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 4ebd10 <hypre_CSRMatrixGetLoadBalancedPartitionEnd> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %R14,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R14,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R14,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JGE 4e92ff <hypre_CSRMatrixTranspose.extracted+0x8f> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x265930(%RIP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x52769c,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x529f11,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4f8100 <hypre_fprintf> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV $0x529e8a,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1d7,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4faac0 <hypre_error_handler> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x38(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x18(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JS 4e930d <hypre_CSRMatrixTranspose.extracted+0x9d> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP %R12,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JLE 4e933f <hypre_CSRMatrixTranspose.extracted+0xcf> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x2658ec(%RIP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x52769c,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x529f20,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4f8100 <hypre_fprintf> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV $0x529e8a,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1d8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4faac0 <hypre_error_handler> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x38(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%RBP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JS 4e9351 <hypre_CSRMatrixTranspose.extracted+0xe1> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP %R12,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JLE 4e9383 <hypre_CSRMatrixTranspose.extracted+0x113> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x2658a8(%RIP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x52769c,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x529f43,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4f8100 <hypre_fprintf> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV $0x529e8a,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1d9,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4faac0 <hypre_error_handler> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x48(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %RAX,%R12 | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA (%R14,%R12,8),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA (,%RAX,8),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %ESI,%ESI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 5011c0 <_intel_fast_memset> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x38(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x10(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RBX,%RAX,8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP (%RBX,%RSI,8),%RAX | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JGE 4e93d4 <hypre_CSRMatrixTranspose.extracted+0x164> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74dab0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 410130 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA 0x1(%RAX),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x20(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %R11,%RCX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA 0x1(%R12),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RCX,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CMP %RCX,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 4e94b1 <hypre_CSRMatrixTranspose.extracted+0x241> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%R11),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP $0x4,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e961c <hypre_CSRMatrixTranspose.extracted+0x3ac> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R8,%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $-0x4,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
CMP %R8,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e94b1 <hypre_CSRMatrixTranspose.extracted+0x241> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA (%R12,%R9,1),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOT %R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %R11,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JMP 4e9456 <hypre_CSRMatrixTranspose.extracted+0x1e6> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R15,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x30(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%R15),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74dad0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 410130 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV (%R15),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x68(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74daf0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4102a0 <__kmpc_masked@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
CMP $0x1,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e95ce <hypre_CSRMatrixTranspose.extracted+0x35e> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP $0x1,%R13 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JLE 4e95c1 <hypre_CSRMatrixTranspose.extracted+0x351> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%R13),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,-0x50(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CMP $0x4,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e9a59 <hypre_CSRMatrixTranspose.extracted+0x7e9> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x50(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $-0x4,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
CMP %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e95c1 <hypre_CSRMatrixTranspose.extracted+0x351> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA 0x1(%RCX),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
IMUL %R9,%RSI | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
DEC %RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
ADD $0x2,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
IMUL %RAX,%RCX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
JMP 4e9570 <hypre_CSRMatrixTranspose.extracted+0x300> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV (%R11),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74db10,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 4100c0 <__kmpc_end_masked@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74db30,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 410130 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
CMPQ $0,-0x48(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
MOV 0x20(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JLE 4e97b0 <hypre_CSRMatrixTranspose.extracted+0x540> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%R12),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
OR %R13,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
SHR $0x20,%RCX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
JE 4e9734 <hypre_CSRMatrixTranspose.extracted+0x4c4> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CQTO | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
IDIV %R13 | 5 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 10 |
MOV -0x58(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RCX,%R12 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JL 4e9745 <hypre_CSRMatrixTranspose.extracted+0x4d5> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e97b0 <hypre_CSRMatrixTranspose.extracted+0x540> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV %R8,%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SHR $0x2,%R9 | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %R12,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
JMP 4e964c <hypre_CSRMatrixTranspose.extracted+0x3dc> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
DIV %R13D | 4 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 6 |
MOV -0x58(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RCX,%R12 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 4e97b0 <hypre_CSRMatrixTranspose.extracted+0x540> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
IMUL %R9,%RDX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %RAX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV (%R14,%RDX,8),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA -0x1(%R9),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP $0x8,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e9847 <hypre_CSRMatrixTranspose.extracted+0x5d7> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %RDI,%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $-0x8,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
CMP %RDI,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e97b0 <hypre_CSRMatrixTranspose.extracted+0x540> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA (%R12,%R8,1),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOT %R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %R9,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JMP 4e979b <hypre_CSRMatrixTranspose.extracted+0x52b> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x30(%RBP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74db50,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 410130 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
CMPQ $0,-0x70(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JE 4e980b <hypre_CSRMatrixTranspose.extracted+0x59b> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x40(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x38(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R10,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x10(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JLE 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x28(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RBX,%RSI,8),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %ESI,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %R10D,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA 0x1(%R10),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
TEST $0x1,%CL | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JNE 4e99ec <hypre_CSRMatrixTranspose.extracted+0x77c> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %RSI,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %R8,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9cf5 <hypre_CSRMatrixTranspose.extracted+0xa85> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV -0x40(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x38(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R9,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x10(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JLE 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV (%RBX,%R11,8),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R11D,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %R9D,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA 0x1(%R9),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
TEST $0x1,%AL | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JNE 4e9be3 <hypre_CSRMatrixTranspose.extracted+0x973> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R11,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RDX,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9c49 <hypre_CSRMatrixTranspose.extracted+0x9d9> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV %RDI,%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SHR $0x3,%R8 | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
LEA 0x7(%R12),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 4e9880 <hypre_CSRMatrixTranspose.extracted+0x610> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x1(%RSI),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x8(%RBX,%RSI,8),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %RDI,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JLE 4e9cca <hypre_CSRMatrixTranspose.extracted+0xa5a> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R8,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x38(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RDI,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x30(%RBP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R8,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9cf5 <hypre_CSRMatrixTranspose.extracted+0xa85> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV -0x50(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
SHR $0x2,%RDX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (,%RAX,4),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RCX,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
LEA (%RAX,%RAX,4),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
LEA -0x1(,%RAX,4),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA (%RAX,%RAX,2),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
DEC %RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
LEA -0x1(,%RAX,2),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x1(%RAX),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 4e9ad4 <hypre_CSRMatrixTranspose.extracted+0x864> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x1(%R11),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x8(%RBX,%R11,8),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %RSI,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JLE 4e9c26 <hypre_CSRMatrixTranspose.extracted+0x9b6> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RSI,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RDX,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9c49 <hypre_CSRMatrixTranspose.extracted+0x9d9> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RDI,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x38(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R8,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9cf5 <hypre_CSRMatrixTranspose.extracted+0xa85> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
ADD $0x58,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
RET | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 0 | 2.13 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
Source file and lines | csr_matop.c:380-560 |
Module | exec |
nb instructions | 300 |
nb uops | 324 |
loop length | 1304 |
used x86 registers | 16 |
used mmx registers | 0 |
used xmm registers | 0 |
used ymm registers | 0 |
used zmm registers | 0 |
nb stack references | 15 |
micro-operation queue | 54.17 cycles |
front end | 54.17 cycles |
P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | |
---|---|---|---|---|---|---|---|---|---|---|---|---|
uops | 22.10 | 22.00 | 20.67 | 20.67 | 16.50 | 22.00 | 21.90 | 16.50 | 16.50 | 16.50 | 22.00 | 20.67 |
cycles | 22.10 | 23.40 | 20.67 | 20.67 | 16.50 | 22.00 | 21.90 | 16.50 | 16.50 | 16.50 | 22.00 | 20.67 |
Cycles executing div or sqrt instructions | 16.00 |
FE+BE cycles | 51.00-51.06 |
Stall cycles | 0.00 |
Front-end | 54.17 |
Dispatch | 23.40 |
DIV/SQRT | 16.00 |
Overall L1 | 54.17 |
all | 0% |
load | 0% |
store | 0% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 0% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 0% |
other | 0% |
all | 11% |
load | 12% |
store | 12% |
mul | NA (no mul vectorizable/vectorized instructions) |
add-sub | 12% |
fma | NA (no fma vectorizable/vectorized instructions) |
div/sqrt | 9% |
other | 10% |
Instruction | Nb FU | P0 | P1 | P2 | P3 | P4 | P5 | P6 | P7 | P8 | P9 | P10 | P11 | Latency | Recip. throughput |
---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|---|
PUSH %RBP | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
MOV %RSP,%RBP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
PUSH %R15 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R14 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R13 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %R12 | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
PUSH %RBX | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 5-12 | 0.50 |
SUB $0x58,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R9,%RBX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R8,%R15 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %RCX,-0x70(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RDX,%R12 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %RDI,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CALL 4f9c80 <hypre_NumActiveThreads> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%R13 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 4f9c90 <hypre_GetThreadNum> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,-0x48(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %R12,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 4ebc90 <hypre_CSRMatrixGetLoadBalancedPartitionBegin> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%R14 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R12,%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CALL 4ebd10 <hypre_CSRMatrixGetLoadBalancedPartitionEnd> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV %RAX,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %R14,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV %R14,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV %R14,-0x40(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV %RCX,-0x38(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
JGE 4e92ff <hypre_CSRMatrixTranspose.extracted+0x8f> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x265930(%RIP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x52769c,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x529f11,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4f8100 <hypre_fprintf> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV $0x529e8a,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1d7,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4faac0 <hypre_error_handler> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x38(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x18(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %R14,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JS 4e930d <hypre_CSRMatrixTranspose.extracted+0x9d> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP %R12,%R14 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JLE 4e933f <hypre_CSRMatrixTranspose.extracted+0xcf> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x2658ec(%RIP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x52769c,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x529f20,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4f8100 <hypre_fprintf> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV $0x529e8a,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1d8,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4faac0 <hypre_error_handler> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x38(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x38(%RBP),%R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
TEST %RCX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JS 4e9351 <hypre_CSRMatrixTranspose.extracted+0xe1> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP %R12,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JLE 4e9383 <hypre_CSRMatrixTranspose.extracted+0x113> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x2658a8(%RIP),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x52769c,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x529f43,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EAX,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4f8100 <hypre_fprintf> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV $0x529e8a,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1d9,%ESI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV $0x1,%EDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %ECX,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4faac0 <hypre_error_handler> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x48(%RBP),%R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %RAX,%R12 | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA (%R14,%R12,8),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA (,%RAX,8),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %ESI,%ESI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 5011c0 <_intel_fast_memset> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x38(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV 0x10(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x40(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RBX,%RAX,8),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP (%RBX,%RSI,8),%RAX | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JGE 4e93d4 <hypre_CSRMatrixTranspose.extracted+0x164> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74dab0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 410130 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x48(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA 0x1(%RAX),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x20(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
IMUL %R11,%RCX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
LEA 0x1(%R12),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RCX,-0x58(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CMP %RCX,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 4e94b1 <hypre_CSRMatrixTranspose.extracted+0x241> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%R11),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP $0x4,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e961c <hypre_CSRMatrixTranspose.extracted+0x3ac> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R8,%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $-0x4,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
CMP %R8,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e94b1 <hypre_CSRMatrixTranspose.extracted+0x241> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA (%R12,%R9,1),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOT %R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %R11,%R9 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JMP 4e9456 <hypre_CSRMatrixTranspose.extracted+0x1e6> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %R15,-0x68(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x30(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%R15),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74dad0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 410130 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV (%R15),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x68(%RBP),%R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74daf0,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CALL 4102a0 <__kmpc_masked@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
CMP $0x1,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e95ce <hypre_CSRMatrixTranspose.extracted+0x35e> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CMP $0x1,%R13 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV -0x30(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JLE 4e95c1 <hypre_CSRMatrixTranspose.extracted+0x351> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%R13),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,-0x50(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
CMP $0x4,%RAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e9a59 <hypre_CSRMatrixTranspose.extracted+0x7e9> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x50(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $-0x4,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
CMP %RAX,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e95c1 <hypre_CSRMatrixTranspose.extracted+0x351> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA 0x1(%RCX),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %RAX,%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
IMUL %R9,%RSI | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
DEC %RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
ADD $0x2,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
IMUL %RAX,%RCX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
JMP 4e9570 <hypre_CSRMatrixTranspose.extracted+0x300> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV (%R11),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74db10,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 4100c0 <__kmpc_end_masked@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74db30,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 410130 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
CMPQ $0,-0x48(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
MOV 0x20(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JLE 4e97b0 <hypre_CSRMatrixTranspose.extracted+0x540> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA -0x1(%R12),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RAX,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
OR %R13,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
SHR $0x20,%RCX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
JE 4e9734 <hypre_CSRMatrixTranspose.extracted+0x4c4> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
CQTO | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 1 | 0.50 |
IDIV %R13 | 5 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 10 |
MOV -0x58(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RCX,%R12 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JL 4e9745 <hypre_CSRMatrixTranspose.extracted+0x4d5> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e97b0 <hypre_CSRMatrixTranspose.extracted+0x540> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV %R8,%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SHR $0x2,%R9 | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV %R12,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
JMP 4e964c <hypre_CSRMatrixTranspose.extracted+0x3dc> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
XOR %EDX,%EDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
DIV %R13D | 4 | 0 | 3 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 11-16 | 6 |
MOV -0x58(%RBP),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RCX,%R12 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JGE 4e97b0 <hypre_CSRMatrixTranspose.extracted+0x540> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
IMUL %R9,%RDX | 1 | 0 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 3 | 1 |
ADD %RAX,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV (%R14,%RDX,8),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA -0x1(%R9),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
CMP $0x8,%RDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e9847 <hypre_CSRMatrixTranspose.extracted+0x5d7> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %RDI,%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
AND $-0x8,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1-2 | 0.20 |
CMP %RDI,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JAE 4e97b0 <hypre_CSRMatrixTranspose.extracted+0x540> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
LEA (%R12,%R8,1),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOT %R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
ADD %R9,%R8 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JMP 4e979b <hypre_CSRMatrixTranspose.extracted+0x52b> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV 0x30(%RBP),%R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x30(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RAX),%ESI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV $0x74db50,%EDI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
CALL 410130 <__kmpc_barrier@plt> | 2 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 0 | 1 |
CMPQ $0,-0x70(%RBP) | 1 | 0.20 | 0.20 | 0.33 | 0.33 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.33 | 1 | 0.33 |
JE 4e980b <hypre_CSRMatrixTranspose.extracted+0x59b> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV -0x40(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x38(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R10,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x10(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JLE 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV 0x28(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV (%RBX,%RSI,8),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %ESI,%ECX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %R10D,%ECX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA 0x1(%R10),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
TEST $0x1,%CL | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JNE 4e99ec <hypre_CSRMatrixTranspose.extracted+0x77c> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %RSI,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %R8,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9cf5 <hypre_CSRMatrixTranspose.extracted+0xa85> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV -0x40(%RBP),%R9 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV -0x38(%RBP),%R11 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R9,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
MOV 0x10(%RBP),%R10 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
JLE 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV (%RBX,%R11,8),%RCX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
MOV %R11D,%EAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SUB %R9D,%EAX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
LEA 0x1(%R9),%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
TEST $0x1,%AL | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 2 | 0.20 |
JNE 4e9be3 <hypre_CSRMatrixTranspose.extracted+0x973> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R11,%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RDX,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9c49 <hypre_CSRMatrixTranspose.extracted+0x9d9> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV %RDI,%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
SHR $0x3,%R8 | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
LEA 0x7(%R12),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 4e9880 <hypre_CSRMatrixTranspose.extracted+0x610> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x1(%RSI),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x8(%RBX,%RSI,8),%RDI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %RDI,%RDX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JLE 4e9cca <hypre_CSRMatrixTranspose.extracted+0xa5a> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
MOV %R8,-0x30(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
MOV -0x38(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RDI,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x30(%RBP),%R8 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R8,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9cf5 <hypre_CSRMatrixTranspose.extracted+0xa85> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
MOV -0x50(%RBP),%RDX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
SHR $0x2,%RDX | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0-2 | 0.50 |
MOV 0x20(%RBP),%RAX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
LEA (,%RAX,4),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RCX,-0x78(%RBP) | 1 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0.50 | 0.50 | 0.50 | 0 | 0 | 1 | 0.50 |
LEA (%RAX,%RAX,4),%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
DEC %RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
LEA -0x1(,%RAX,4),%RSI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA (%RAX,%RAX,2),%RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
DEC %RDI | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
LEA -0x1(,%RAX,2),%R8 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x1(%RAX),%R9 | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
JMP 4e9ad4 <hypre_CSRMatrixTranspose.extracted+0x864> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 5.84 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPL (%RAX) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
LEA -0x1(%R11),%RAX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV -0x8(%RBX,%R11,8),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %RSI,%RCX | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JLE 4e9c26 <hypre_CSRMatrixTranspose.extracted+0x9b6> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RSI,%RCX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
CMP %RDX,%R11 | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9c49 <hypre_CSRMatrixTranspose.extracted+0x9d9> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
JMP 4e9cd6 <hypre_CSRMatrixTranspose.extracted+0xa66> | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 2.08 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
MOV %RDI,%RDX | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
MOV -0x38(%RBP),%RSI | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1 | 0.33 |
CMP %R8,%RSI | 1 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0.20 | 0 | 0 | 0 | 0.20 | 0 | 1 | 0.20 |
JNE 4e9cf5 <hypre_CSRMatrixTranspose.extracted+0xa85> | 1 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0 | 0 | 0.50 |
ADD $0x58,%RSP | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 1 | 0.17 |
POP %RBX | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R12 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R13 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R14 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %R15 | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
POP %RBP | 1 | 0 | 0 | 0.33 | 0.33 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.33 | 1-6 | 0.33 |
RET | 1 | 0.50 | 0 | 0.33 | 0.33 | 0 | 0 | 0.50 | 0 | 0 | 0 | 0 | 0.33 | 0 | 2.13 |
NOPW %CS:(%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
NOPW (%RAX,%RAX,1) | 1 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0 | 0.17 |
Name | Coverage (%) | Time (s) |
---|---|---|
▼hypre_CSRMatrixTranspose.extracted– | 0.17 | 0.04 |
▼Loop 3434 - csr_matop.c:540-548 - exec– | 0.03 | 0.01 |
○Loop 3435 - csr_matop.c:541-548 - exec | 0.05 | 0.01 |
○Loop 3436 - csr_matop.c:540-548 - exec | 0.05 | 0.01 |
○Loop 3439 - csr_matop.c:380-527 - exec | 0.01 | 0.01 |
○Loop 3444 - csr_matop.c:483-485 - exec | 0.01 | 0.01 |
○Loop 3443 - csr_matop.c:380-500 - exec | 0.01 | 0.01 |
▼Loop 3430 - csr_matop.c:553-560 - exec– | 0 | 0 |
○Loop 3432 - csr_matop.c:553-560 - exec | 0 | 0 |
○Loop 3431 - csr_matop.c:554-560 - exec | 0 | 0 |
○Loop 3433 - csr_matop.c:554-560 - exec | 0 | 0 |
○Loop 3441 - csr_matop.c:380-513 - exec | 0 | 0 |
○Loop 3442 - csr_matop.c:380-500 - exec | 0 | 0 |
○Loop 3438 - csr_matop.c:380-527 - exec | 0 | 0 |
○Loop 3437 - csr_matop.c:541-548 - exec | 0 | 0 |
○Loop 3440 - csr_matop.c:380-513 - exec | 0 | 0 |