/external/XNNPACK/src/f32-gemm/gen/ |
D | 1x16s4-minmax-fma3-broadcast.c | 60 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemm_minmax_ukernel_1x16s4__fma3_broadcast() local 100 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemm_minmax_ukernel_1x16s4__fma3_broadcast() local
|
D | 3x16s4-minmax-fma3-broadcast.c | 86 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemm_minmax_ukernel_3x16s4__fma3_broadcast() local 152 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemm_minmax_ukernel_3x16s4__fma3_broadcast() local
|
D | 4x16s4-minmax-fma3-broadcast.c | 99 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemm_minmax_ukernel_4x16s4__fma3_broadcast() local 178 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemm_minmax_ukernel_4x16s4__fma3_broadcast() local
|
/external/XNNPACK/src/f32-gemm/gen-inc/ |
D | 1x16s4inc-minmax-fma3-broadcast.c | 62 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemminc_minmax_ukernel_1x16s4__fma3_broadcast() local 102 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemminc_minmax_ukernel_1x16s4__fma3_broadcast() local
|
D | 3x16s4inc-minmax-fma3-broadcast.c | 88 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemminc_minmax_ukernel_3x16s4__fma3_broadcast() local 154 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemminc_minmax_ukernel_3x16s4__fma3_broadcast() local
|
D | 4x16s4inc-minmax-fma3-broadcast.c | 101 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemminc_minmax_ukernel_4x16s4__fma3_broadcast() local 180 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_gemminc_minmax_ukernel_4x16s4__fma3_broadcast() local
|
/external/XNNPACK/src/f32-igemm/gen/ |
D | 1x16s4-minmax-fma3-broadcast.c | 73 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_igemm_minmax_ukernel_1x16s4__fma3_broadcast() local 113 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_igemm_minmax_ukernel_1x16s4__fma3_broadcast() local
|
D | 3x16s4-minmax-fma3-broadcast.c | 105 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_igemm_minmax_ukernel_3x16s4__fma3_broadcast() local 171 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_igemm_minmax_ukernel_3x16s4__fma3_broadcast() local
|
D | 4x16s4-minmax-fma3-broadcast.c | 121 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_igemm_minmax_ukernel_4x16s4__fma3_broadcast() local 200 const __m256 vb01234567c1 = _mm256_load_ps(w + 16); in xnn_f32_igemm_minmax_ukernel_4x16s4__fma3_broadcast() local
|
/external/XNNPACK/src/qs8-gemm/gen/ |
D | 1x8-minmax-rndnu-neon-mull-addw-dup.c | 55 … const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((uintptr_t) w + 8 * sizeof(int8_t)); in xnn_qs8_gemm_minmax_rndnu_ukernel_1x8__neon_mull_addw_dup() local 103 … const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((uintptr_t) w + 8 * sizeof(int8_t)); in xnn_qs8_gemm_minmax_rndnu_ukernel_1x8__neon_mull_addw_dup() local
|
D | 1x8-minmax-rndnu-neon-mlal-lane.c | 56 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qs8_gemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane() local 107 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qs8_gemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane() local
|
D | 1x8-minmax-rndnu-neon-mlal-lane-prfm.c | 56 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qs8_gemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane_prfm() local 108 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qs8_gemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane_prfm() local
|
/external/XNNPACK/src/qs8-igemm/gen/ |
D | 1x8-minmax-rndnu-neon-mlal-lane-prfm.c | 67 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qs8_igemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane_prfm() local 119 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qs8_igemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane_prfm() local
|
D | 1x8-minmax-rndnu-neon-mull-addw-dup.c | 66 … const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((uintptr_t) w + 8 * sizeof(int8_t)); in xnn_qs8_igemm_minmax_rndnu_ukernel_1x8__neon_mull_addw_dup() local 114 … const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((uintptr_t) w + 8 * sizeof(int8_t)); in xnn_qs8_igemm_minmax_rndnu_ukernel_1x8__neon_mull_addw_dup() local
|
D | 1x8-minmax-rndnu-neon-mlal-lane.c | 67 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qs8_igemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane() local 118 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qs8_igemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane() local
|
/external/XNNPACK/src/qc8-gemm/gen/ |
D | 1x8-minmax-fp32-neonv8-mlal-lane-prfm.c | 57 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_gemm_minmax_fp32_ukernel_1x8__neonv8_mlal_lane_prfm() local 109 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_gemm_minmax_fp32_ukernel_1x8__neonv8_mlal_lane_prfm() local
|
D | 1x8-minmax-fp32-neonv8-mlal-lane.c | 57 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_gemm_minmax_fp32_ukernel_1x8__neonv8_mlal_lane() local 108 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_gemm_minmax_fp32_ukernel_1x8__neonv8_mlal_lane() local
|
D | 1x8-minmax-fp32-neon-mlal-lane.c | 56 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_gemm_minmax_fp32_ukernel_1x8__neon_mlal_lane() local 107 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_gemm_minmax_fp32_ukernel_1x8__neon_mlal_lane() local
|
D | 1x8-minmax-fp32-neon-mlal-lane-prfm.c | 56 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_gemm_minmax_fp32_ukernel_1x8__neon_mlal_lane_prfm() local 108 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_gemm_minmax_fp32_ukernel_1x8__neon_mlal_lane_prfm() local
|
/external/XNNPACK/src/qu8-gemm/gen/ |
D | 1x8-minmax-rndnu-neon-mlal-lane.c | 57 const uint8x8_t vb01234567c1 = vld1_u8(w); w = (const void*) ((const uint8_t*) w + 8); in xnn_qu8_gemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane() local 108 const uint8x8_t vb01234567c1 = vld1_u8(w); w = (const void*) ((const uint8_t*) w + 8); in xnn_qu8_gemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane() local
|
D | 1x8-minmax-fp32-neon-mlal-lane.c | 57 const uint8x8_t vb01234567c1 = vld1_u8(w); w = (const void*) ((const uint8_t*) w + 8); in xnn_qu8_gemm_minmax_fp32_ukernel_1x8__neon_mlal_lane() local 108 const uint8x8_t vb01234567c1 = vld1_u8(w); w = (const void*) ((const uint8_t*) w + 8); in xnn_qu8_gemm_minmax_fp32_ukernel_1x8__neon_mlal_lane() local
|
/external/XNNPACK/src/qu8-igemm/gen/ |
D | 1x8-minmax-rndnu-neon-mlal-lane.c | 68 const uint8x8_t vb01234567c1 = vld1_u8(w); w = (const void*) ((const uint8_t*) w + 8); in xnn_qu8_igemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane() local 119 const uint8x8_t vb01234567c1 = vld1_u8(w); w = (const void*) ((const uint8_t*) w + 8); in xnn_qu8_igemm_minmax_rndnu_ukernel_1x8__neon_mlal_lane() local
|
D | 1x8-minmax-fp32-neon-mlal-lane.c | 68 const uint8x8_t vb01234567c1 = vld1_u8(w); w = (const void*) ((const uint8_t*) w + 8); in xnn_qu8_igemm_minmax_fp32_ukernel_1x8__neon_mlal_lane() local 119 const uint8x8_t vb01234567c1 = vld1_u8(w); w = (const void*) ((const uint8_t*) w + 8); in xnn_qu8_igemm_minmax_fp32_ukernel_1x8__neon_mlal_lane() local
|
/external/XNNPACK/src/qc8-igemm/gen/ |
D | 1x8-minmax-fp32-neonv8-mlal-lane.c | 68 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_igemm_minmax_fp32_ukernel_1x8__neonv8_mlal_lane() local 119 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_igemm_minmax_fp32_ukernel_1x8__neonv8_mlal_lane() local
|
D | 1x8-minmax-fp32-neon-mlal-lane-prfm.c | 67 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_igemm_minmax_fp32_ukernel_1x8__neon_mlal_lane_prfm() local 119 const int8x8_t vb01234567c1 = vld1_s8(w); w = (const void*) ((const int8_t*) w + 8); in xnn_qc8_igemm_minmax_fp32_ukernel_1x8__neon_mlal_lane_prfm() local
|