Lines Matching refs:vi
69 float32x4_t vi${M}x1357 = vmovq_n_f32(0.0f);
77 const float32x4x2_t vi${M}x8ACE9BDF = vld2q_f32(i${M}); i${M} += 8;
81 float32x4_t vo${M}p1 = vmulq_lane_f32(vi${2*M}x8ACE9BDF.val[0], vget_high_f32(vw0123), 0);
83 … vo${M}p0 = ${VMULADDQ_LANE_F32}(vo${M}p0, vi${2*M}x8ACE9BDF.val[0], vget_high_f32(vw0123), 0);
87 … float32x4_t vo${M}p2 = vmulq_lane_f32(vi${2*M+1}x8ACE9BDF.val[0], vget_low_f32(vw4567), 1);
89 … vo${M}p0 = ${VMULADDQ_LANE_F32}(vo${M}p0, vi${2*M+1}x8ACE9BDF.val[0], vget_low_f32(vw4567), 1);
93 float32x4_t vo${M}p3 = vmulq_lane_f32(vi${2*M+2}x8ACE9BDF.val[0], vw89, 0);
95 …vo${M}p${4 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${4 % ACCUMULATORS}, vi${2*M+2}x8ACE9BDF.…
98 const float32x4_t vi${M}x7BDF = vextq_f32(vi${M}x1357, vi${M}x8ACE9BDF.val[1], 3);
99 vi${M}x1357 = vi${M}x8ACE9BDF.val[1];
103 float32x4_t vo${M}p4 = vmulq_lane_f32(vi${2*M}x7BDF, vget_low_f32(vw0123), 1);
105 …vo${M}p${5 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${5 % ACCUMULATORS}, vi${2*M}x7BDF, vget_…
109 float32x4_t vo${M}p5 = vmulq_lane_f32(vi${2*M+1}x7BDF, vget_low_f32(vw4567), 0);
111 …vo${M}p${6 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${6 % ACCUMULATORS}, vi${2*M+1}x7BDF, vge…
115 float32x4_t vo${M}p6 = vmulq_lane_f32(vi${2*M+2}x7BDF, vget_low_f32(vw4567), 1);
117 …vo${M}p${7 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${7 % ACCUMULATORS}, vi${2*M+2}x7BDF, vge…
120 …vo${M}p${8 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${8 % ACCUMULATORS}, vi${2*M}x8ACE9BDF.va…
123 …vo${M}p${9 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${9 % ACCUMULATORS}, vi${2*M+1}x8ACE9BDF.…
126 …vo${M}p${10 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${10 % ACCUMULATORS}, vi${2*M+2}x8ACE9BD…
153 const float32x4x2_t vi${M}x8ACE9BDF = vld2q_f32(i${M});
156 …const float32x4_t vi${M}x8ACE = vreinterpretq_f32_u32(vandq_u32(vmask_even, vreinterpretq_u32_f32(…
157 …const float32x4_t vi${M}x9BDF = vreinterpretq_f32_u32(vandq_u32(vmask_odd, vreinterpretq_u32_f32(…
161 float32x4_t vo${M}p1 = vmulq_lane_f32(vi${2*M}x8ACE, vget_high_f32(vw0123), 0);
163 vo${M}p0 = ${VMULADDQ_LANE_F32}(vo${M}p0, vi${2*M}x8ACE, vget_high_f32(vw0123), 0);
167 float32x4_t vo${M}p2 = vmulq_lane_f32(vi${2*M+1}x8ACE, vget_low_f32(vw4567), 1);
169 vo${M}p0 = ${VMULADDQ_LANE_F32}(vo${M}p0, vi${2*M+1}x8ACE, vget_low_f32(vw4567), 1);
173 float32x4_t vo${M}p3 = vmulq_lane_f32(vi${2*M+2}x8ACE, vw89, 0);
175 …vo${M}p${4 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${4 % ACCUMULATORS}, vi${2*M+2}x8ACE, vw8…
178 const float32x4_t vi${M}x7BDF = vextq_f32(vi${M}x1357, vi${M}x9BDF, 3);
182 float32x4_t vo${M}p4 = vmulq_lane_f32(vi${2*M}x7BDF, vget_low_f32(vw0123), 1);
184 …vo${M}p${5 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${5 % ACCUMULATORS}, vi${2*M}x7BDF, vget_…
188 float32x4_t vo${M}p5 = vmulq_lane_f32(vi${2*M+1}x7BDF, vget_low_f32(vw4567), 0);
190 …vo${M}p${6 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${6 % ACCUMULATORS}, vi${2*M+1}x7BDF, vge…
194 float32x4_t vo${M}p6 = vmulq_lane_f32(vi${2*M+2}x7BDF, vget_low_f32(vw4567), 1);
196 …vo${M}p${7 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${7 % ACCUMULATORS}, vi${2*M+2}x7BDF, vge…
199 …vo${M}p${8 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${8 % ACCUMULATORS}, vi${2*M}x9BDF, vget_…
202 …vo${M}p${9 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${9 % ACCUMULATORS}, vi${2*M+1}x9BDF, vge…
205 …vo${M}p${10 % ACCUMULATORS} = ${VMULADDQ_LANE_F32}(vo${M}p${10 % ACCUMULATORS}, vi${2*M+2}x9BDF, v…