/external/XNNPACK/src/qs8-gemm/gen/ |
D | 2x8c4-minmax-rndnu-neon-mlal-ld2r.c | 96 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld2r() local 171 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld2r() local 224 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld2r() local
|
D | 2x8c4-minmax-rndnu-neon-mlal-dup.c | 96 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_dup() local 171 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_dup() local 224 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_dup() local
|
D | 2x8c4-minmax-fp32-neon-mlal-ld2r.c | 96 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local 171 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local 224 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local
|
D | 2x8c4-minmax-fp32-neon-mlal-dup.c | 96 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local 171 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local 224 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-dup.c | 97 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local 172 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local 225 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-ld2r.c | 97 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local 172 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local 225 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-ld1r.c | 101 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld1r() local 178 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld1r() local 231 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld1r() local
|
D | 2x8c4-minmax-rndnu-neon-mlal-ld1r.c | 100 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld1r() local 177 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld1r() local 230 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld1r() local
|
D | 2x8c4-minmax-fp32-neon-mlal-ld1r.c | 100 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld1r() local 177 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld1r() local 230 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld1r() local
|
/external/XNNPACK/src/qc8-gemm/gen/ |
D | 2x8c4-minmax-fp32-neon-mlal-ld2r.c | 96 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local 171 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local 224 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-ld2r.c | 97 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local 172 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local 225 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-dup.c | 97 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local 172 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local 225 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-ld1r.c | 101 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld1r() local 178 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld1r() local 231 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld1r() local
|
D | 2x8c4-minmax-fp32-neon-mlal-dup.c | 96 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local 171 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local 224 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_gemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local
|
/external/XNNPACK/src/qs8-igemm/gen/ |
D | 2x8c4-minmax-fp32-neon-mlal-ld2r.c | 109 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local 184 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local 237 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local
|
D | 2x8c4-minmax-rndnu-neon-mlal-ld1r.c | 113 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld1r() local 190 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld1r() local 243 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld1r() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-dup.c | 110 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local 185 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local 238 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local
|
D | 2x8c4-minmax-rndnu-neon-mlal-dup.c | 109 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_dup() local 184 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_dup() local 237 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_dup() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-ld2r.c | 110 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local 185 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local 238 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local
|
D | 2x8c4-minmax-fp32-neon-mlal-dup.c | 109 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local 184 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local 237 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local
|
D | 2x8c4-minmax-rndnu-neon-mlal-ld2r.c | 109 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld2r() local 184 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld2r() local 237 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qs8_igemm_minmax_rndnu_ukernel_2x8c4__neon_mlal_ld2r() local
|
/external/XNNPACK/src/qc8-igemm/gen/ |
D | 2x8c4-minmax-fp32-neon-mlal-dup.c | 109 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local 184 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local 237 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_dup() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-dup.c | 110 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local 185 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local 238 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_dup() local
|
D | 2x8c4-minmax-fp32-neonv8-mlal-ld2r.c | 110 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local 185 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local 238 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neonv8_mlal_ld2r() local
|
D | 2x8c4-minmax-fp32-neon-mlal-ld2r.c | 109 int16x8_t vprod1x45c0 = vmull_s8(vb45c0x0, va1c0x0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local 184 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local 237 const int16x8_t vprod1x45c0 = vmull_s8(vb45c0, va1c0); in xnn_qc8_igemm_minmax_fp32_ukernel_2x8c4__neon_mlal_ld2r() local
|