Coverage Report

Created: 2026-08-31 06:19

next uncovered line (L), next uncovered region (R), next uncovered branch (B)
/work/dav1d/src/mc_tmpl.c
Line
Count
Source
1
/*
2
 * Copyright © 2018, VideoLAN and dav1d authors
3
 * Copyright © 2018, Two Orioles, LLC
4
 * All rights reserved.
5
 *
6
 * Redistribution and use in source and binary forms, with or without
7
 * modification, are permitted provided that the following conditions are met:
8
 *
9
 * 1. Redistributions of source code must retain the above copyright notice, this
10
 *    list of conditions and the following disclaimer.
11
 *
12
 * 2. Redistributions in binary form must reproduce the above copyright notice,
13
 *    this list of conditions and the following disclaimer in the documentation
14
 *    and/or other materials provided with the distribution.
15
 *
16
 * THIS SOFTWARE IS PROVIDED BY THE COPYRIGHT HOLDERS AND CONTRIBUTORS "AS IS" AND
17
 * ANY EXPRESS OR IMPLIED WARRANTIES, INCLUDING, BUT NOT LIMITED TO, THE IMPLIED
18
 * WARRANTIES OF MERCHANTABILITY AND FITNESS FOR A PARTICULAR PURPOSE ARE
19
 * DISCLAIMED. IN NO EVENT SHALL THE COPYRIGHT OWNER OR CONTRIBUTORS BE LIABLE FOR
20
 * ANY DIRECT, INDIRECT, INCIDENTAL, SPECIAL, EXEMPLARY, OR CONSEQUENTIAL DAMAGES
21
 * (INCLUDING, BUT NOT LIMITED TO, PROCUREMENT OF SUBSTITUTE GOODS OR SERVICES;
22
 * LOSS OF USE, DATA, OR PROFITS; OR BUSINESS INTERRUPTION) HOWEVER CAUSED AND
23
 * ON ANY THEORY OF LIABILITY, WHETHER IN CONTRACT, STRICT LIABILITY, OR TORT
24
 * (INCLUDING NEGLIGENCE OR OTHERWISE) ARISING IN ANY WAY OUT OF THE USE OF THIS
25
 * SOFTWARE, EVEN IF ADVISED OF THE POSSIBILITY OF SUCH DAMAGE.
26
 */
27
28
#include "config.h"
29
30
#include <stdlib.h>
31
#include <string.h>
32
33
#include "common/attributes.h"
34
#include "common/intops.h"
35
36
#include "src/mc.h"
37
#include "src/tables.h"
38
39
#if BITDEPTH == 8
40
1.80M
#define get_intermediate_bits(bitdepth_max) 4
41
// Output in interval [-5132, 9212], fits in int16_t as is
42
120M
#define PREP_BIAS 0
43
#else
44
// 4 for 10 bits/component, 2 for 12 bits/component
45
#define get_intermediate_bits(bitdepth_max) (14 - bitdepth_from_max(bitdepth_max))
46
// Output in interval [-20588, 36956] (10-bit), [-20602, 36983] (12-bit)
47
// Subtract a bias to ensure the output fits in int16_t
48
#define PREP_BIAS 8192
49
#endif
50
51
static NOINLINE void
52
put_c(pixel *dst, const ptrdiff_t dst_stride,
53
      const pixel *src, const ptrdiff_t src_stride, const int w, int h)
54
518k
{
55
7.72M
    do {
56
7.72M
        pixel_copy(dst, src, w);
57
58
7.72M
        dst += dst_stride;
59
7.72M
        src += src_stride;
60
7.72M
    } while (--h);
61
518k
}
62
63
static NOINLINE void
64
prep_c(int16_t *tmp, const pixel *src, const ptrdiff_t src_stride,
65
       const int w, int h HIGHBD_DECL_SUFFIX)
66
87.9k
{
67
87.9k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
68
2.00M
    do {
69
43.4M
        for (int x = 0; x < w; x++)
70
41.4M
            tmp[x] = (src[x] << intermediate_bits) - PREP_BIAS;
71
72
2.00M
        tmp += w;
73
2.00M
        src += src_stride;
74
2.00M
    } while (--h);
75
87.9k
}
76
77
#define FILTER_8TAP(src, x, F, stride) \
78
233M
    (F[0] * src[x + -3 * stride] + \
79
233M
     F[1] * src[x + -2 * stride] + \
80
233M
     F[2] * src[x + -1 * stride] + \
81
233M
     F[3] * src[x + +0 * stride] + \
82
233M
     F[4] * src[x + +1 * stride] + \
83
233M
     F[5] * src[x + +2 * stride] + \
84
233M
     F[6] * src[x + +3 * stride] + \
85
233M
     F[7] * src[x + +4 * stride])
86
87
#define FILTER_8TAP2(src, x, F) \
88
89.1M
    (F[0] * src[0][x] + \
89
89.1M
     F[1] * src[1][x] + \
90
89.1M
     F[2] * src[2][x] + \
91
89.1M
     F[3] * src[3][x] + \
92
89.1M
     F[4] * src[4][x] + \
93
89.1M
     F[5] * src[5][x] + \
94
89.1M
     F[6] * src[6][x] + \
95
89.1M
     F[7] * src[7][x])
96
97
#define DAV1D_FILTER_8TAP_RND(src, x, F, stride, sh) \
98
229M
    ((FILTER_8TAP(src, x, F, stride) + ((1 << (sh)) >> 1)) >> (sh))
99
100
#define DAV1D_FILTER_8TAP_RND2(src, x, F, stride, rnd, sh) \
101
4.22M
    ((FILTER_8TAP(src, x, F, stride) + (rnd)) >> (sh))
102
103
#define DAV1D_FILTER_8TAP_RND3(src, x, F, sh) \
104
89.1M
    ((FILTER_8TAP2(src, x, F) + ((1 << (sh)) >> 1)) >> (sh))
105
106
#define DAV1D_FILTER_8TAP_CLIP(src, x, F, stride, sh) \
107
23.3M
    iclip_pixel(DAV1D_FILTER_8TAP_RND(src, x, F, stride, sh))
108
109
#define DAV1D_FILTER_8TAP_CLIP2(src, x, F, stride, rnd, sh) \
110
4.22M
    iclip_pixel(DAV1D_FILTER_8TAP_RND2(src, x, F, stride, rnd, sh))
111
112
#define DAV1D_FILTER_8TAP_CLIP3(src, x, F, sh) \
113
61.1M
    iclip_pixel(DAV1D_FILTER_8TAP_RND3(src, x, F, sh))
114
115
#define GET_H_FILTER(mx) \
116
161M
    const int8_t *const fh = !(mx) ? NULL : w > 4 ? \
117
149M
        dav1d_mc_subpel_filters[filter_type & 3][(mx) - 1] : \
118
149M
        dav1d_mc_subpel_filters[3 + (filter_type & 1)][(mx) - 1]
119
120
#define GET_V_FILTER(my) \
121
5.23M
    const int8_t *const fv = !(my) ? NULL : h > 4 ? \
122
3.71M
        dav1d_mc_subpel_filters[filter_type >> 2][(my) - 1] : \
123
3.71M
        dav1d_mc_subpel_filters[3 + ((filter_type >> 2) & 1)][(my) - 1]
124
125
#define GET_FILTERS() \
126
717k
    GET_H_FILTER(mx); \
127
717k
    GET_V_FILTER(my)
128
129
static NOINLINE void
130
put_8tap_c(pixel *dst, ptrdiff_t dst_stride,
131
           const pixel *src, ptrdiff_t src_stride,
132
           const int w, int h, const int mx, const int my,
133
           const int filter_type HIGHBD_DECL_SUFFIX)
134
580k
{
135
580k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
136
580k
    const int intermediate_rnd = 32 + ((1 << (6 - intermediate_bits)) >> 1);
137
138
580k
    GET_FILTERS();
139
580k
    dst_stride = PXSTRIDE(dst_stride);
140
580k
    src_stride = PXSTRIDE(src_stride);
141
142
580k
    if (fh) {
143
126k
        if (fv) {
144
90.3k
            int tmp_h = h + 7;
145
90.3k
            int16_t mid[128 * 135], *mid_ptr = mid;
146
147
90.3k
            src -= src_stride * 3;
148
1.57M
            do {
149
23.3M
                for (int x = 0; x < w; x++)
150
21.8M
                    mid_ptr[x] = DAV1D_FILTER_8TAP_RND(src, x, fh, 1,
151
1.57M
                                                       6 - intermediate_bits);
152
153
1.57M
                mid_ptr += 128;
154
1.57M
                src += src_stride;
155
1.57M
            } while (--tmp_h);
156
157
90.3k
            mid_ptr = mid + 128 * 3;
158
931k
            do {
159
16.6M
                for (int x = 0; x < w; x++)
160
15.7M
                    dst[x] = DAV1D_FILTER_8TAP_CLIP(mid_ptr, x, fv, 128,
161
931k
                                                    6 + intermediate_bits);
162
163
931k
                mid_ptr += 128;
164
931k
                dst += dst_stride;
165
931k
            } while (--h);
166
90.3k
        } else {
167
293k
            do {
168
4.51M
                for (int x = 0; x < w; x++) {
169
4.22M
                    dst[x] = DAV1D_FILTER_8TAP_CLIP2(src, x, fh, 1,
170
4.22M
                                                     intermediate_rnd, 6);
171
4.22M
                }
172
173
293k
                dst += dst_stride;
174
293k
                src += src_stride;
175
293k
            } while (--h);
176
35.9k
        }
177
454k
    } else if (fv) {
178
511k
        do {
179
8.18M
            for (int x = 0; x < w; x++)
180
7.66M
                dst[x] = DAV1D_FILTER_8TAP_CLIP(src, x, fv, src_stride, 6);
181
182
511k
            dst += dst_stride;
183
511k
            src += src_stride;
184
511k
        } while (--h);
185
61.4k
    } else
186
393k
        put_c(dst, dst_stride, src, src_stride, w, h);
187
580k
}
188
189
static NOINLINE void
190
put_8tap_scaled_c(pixel *dst, const ptrdiff_t dst_stride,
191
                  const pixel *src, ptrdiff_t src_stride,
192
                  const int w, int h, const int mx, int my,
193
                  const int dx, const int dy, const int filter_type
194
                  HIGHBD_DECL_SUFFIX)
195
268k
{
196
268k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
197
268k
    const int intermediate_rnd = (1 << intermediate_bits) >> 1;
198
268k
    int16_t mid[128 * 8];
199
268k
    int16_t *mid_ptrs[8];
200
268k
    int in_y = -8;
201
268k
    src_stride = PXSTRIDE(src_stride);
202
203
2.41M
    for (int i = 0; i < 8; i++)
204
2.15M
        mid_ptrs[i] = &mid[128 * i];
205
206
268k
    src -= src_stride * 3;
207
208
3.30M
    for (int y = 0; y < h; y++) {
209
3.03M
        int x;
210
3.03M
        int src_y = my >> 10;
211
3.03M
        GET_V_FILTER((my & 0x3ff) >> 6);
212
213
7.92M
        while (in_y < src_y) {
214
4.89M
            int imx = mx, ioff = 0;
215
4.89M
            int16_t *mid_ptr = mid_ptrs[0];
216
217
39.1M
            for (int i = 0; i < 7; i++)
218
34.2M
                mid_ptrs[i] = mid_ptrs[i + 1];
219
4.89M
            mid_ptrs[7] = mid_ptr;
220
221
109M
            for (x = 0; x < w; x++) {
222
104M
                GET_H_FILTER(imx >> 6);
223
104M
                mid_ptr[x] = fh ? DAV1D_FILTER_8TAP_RND(src, ioff, fh, 1,
224
104M
                                                        6 - intermediate_bits) :
225
104M
                                  src[ioff] << intermediate_bits;
226
104M
                imx += dx;
227
104M
                ioff += imx >> 10;
228
104M
                imx &= 0x3ff;
229
104M
            }
230
231
4.89M
            src += src_stride;
232
4.89M
            in_y++;
233
4.89M
        }
234
235
74.7M
        for (x = 0; x < w; x++)
236
71.7M
            dst[x] = fv ? DAV1D_FILTER_8TAP_CLIP3(mid_ptrs, x, fv,
237
71.7M
                                                  6 + intermediate_bits) :
238
71.7M
                          iclip_pixel((mid_ptrs[3][x] + intermediate_rnd) >>
239
10.5M
                                              intermediate_bits);
240
241
3.03M
        my += dy;
242
3.03M
        dst += PXSTRIDE(dst_stride);
243
3.03M
    }
244
268k
}
245
246
static NOINLINE void
247
prep_8tap_c(int16_t *tmp, const pixel *src, ptrdiff_t src_stride,
248
            const int w, int h, const int mx, const int my,
249
            const int filter_type HIGHBD_DECL_SUFFIX)
250
136k
{
251
136k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
252
136k
    GET_FILTERS();
253
136k
    src_stride = PXSTRIDE(src_stride);
254
255
136k
    if (fh) {
256
45.1k
        if (fv) {
257
33.0k
            int tmp_h = h + 7;
258
33.0k
            int16_t mid[128 * 135], *mid_ptr = mid;
259
260
33.0k
            src -= src_stride * 3;
261
832k
            do {
262
16.7M
                for (int x = 0; x < w; x++)
263
15.8M
                    mid_ptr[x] = DAV1D_FILTER_8TAP_RND(src, x, fh, 1,
264
832k
                                                       6 - intermediate_bits);
265
266
832k
                mid_ptr += 128;
267
832k
                src += src_stride;
268
832k
            } while (--tmp_h);
269
270
33.0k
            mid_ptr = mid + 128 * 3;
271
603k
            do {
272
13.7M
                for (int x = 0; x < w; x++) {
273
13.1M
                    int t = DAV1D_FILTER_8TAP_RND(mid_ptr, x, fv, 128, 6) -
274
13.1M
                                  PREP_BIAS;
275
13.1M
                    assert(t >= INT16_MIN && t <= INT16_MAX);
276
13.1M
                    tmp[x] = t;
277
13.1M
                }
278
279
603k
                mid_ptr += 128;
280
603k
                tmp += w;
281
603k
            } while (--h);
282
33.0k
        } else {
283
182k
            do {
284
3.41M
                for (int x = 0; x < w; x++)
285
3.23M
                    tmp[x] = DAV1D_FILTER_8TAP_RND(src, x, fh, 1,
286
3.23M
                                                   6 - intermediate_bits) -
287
3.23M
                             PREP_BIAS;
288
289
182k
                tmp += w;
290
182k
                src += src_stride;
291
182k
            } while (--h);
292
12.0k
        }
293
91.6k
    } else if (fv) {
294
122k
        do {
295
2.28M
            for (int x = 0; x < w; x++)
296
2.16M
                tmp[x] = DAV1D_FILTER_8TAP_RND(src, x, fv, src_stride,
297
2.16M
                                               6 - intermediate_bits) -
298
2.16M
                         PREP_BIAS;
299
300
122k
            tmp += w;
301
122k
            src += src_stride;
302
122k
        } while (--h);
303
7.96k
    } else
304
83.6k
        prep_c(tmp, src, src_stride, w, h HIGHBD_TAIL_SUFFIX);
305
136k
}
306
307
static NOINLINE void
308
prep_8tap_scaled_c(int16_t *tmp, const pixel *src, ptrdiff_t src_stride,
309
                   const int w, int h, const int mx, int my,
310
                   const int dx, const int dy, const int filter_type
311
                   HIGHBD_DECL_SUFFIX)
312
88.0k
{
313
88.0k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
314
88.0k
    int16_t mid[128 * 8];
315
88.0k
    int16_t *mid_ptrs[8];
316
88.0k
    int in_y = -8;
317
88.0k
    src_stride = PXSTRIDE(src_stride);
318
319
792k
    for (int i = 0; i < 8; i++)
320
704k
        mid_ptrs[i] = &mid[128 * i];
321
322
88.0k
    src -= src_stride * 3;
323
324
1.57M
    for (int y = 0; y < h; y++) {
325
1.48M
        int x;
326
1.48M
        int src_y = my >> 10;
327
1.48M
        GET_V_FILTER((my & 0x3ff) >> 6);
328
329
3.49M
        while (in_y < src_y) {
330
2.00M
            int imx = mx, ioff = 0;
331
2.00M
            int16_t *mid_ptr = mid_ptrs[0];
332
333
16.0M
            for (int i = 0; i < 7; i++)
334
14.0M
                mid_ptrs[i] = mid_ptrs[i + 1];
335
2.00M
            mid_ptrs[7] = mid_ptr;
336
337
58.6M
            for (x = 0; x < w; x++) {
338
56.6M
                GET_H_FILTER(imx >> 6);
339
56.6M
                mid_ptr[x] = fh ? DAV1D_FILTER_8TAP_RND(src, ioff, fh, 1,
340
56.6M
                                                        6 - intermediate_bits) :
341
56.6M
                                  src[ioff] << intermediate_bits;
342
56.6M
                imx += dx;
343
56.6M
                ioff += imx >> 10;
344
56.6M
                imx &= 0x3ff;
345
56.6M
            }
346
347
2.00M
            src += src_stride;
348
2.00M
            in_y++;
349
2.00M
        }
350
351
49.7M
        for (x = 0; x < w; x++)
352
48.2M
            tmp[x] = (fv ? DAV1D_FILTER_8TAP_RND3(mid_ptrs, x, fv, 6)
353
48.2M
                         : mid_ptrs[3][x]) - PREP_BIAS;
354
355
1.48M
        my += dy;
356
1.48M
        tmp += w;
357
1.48M
    }
358
88.0k
}
359
360
#define filter_fns(type, type_h, type_v) \
361
static void put_8tap_##type##_c(pixel *const dst, \
362
                                const ptrdiff_t dst_stride, \
363
                                const pixel *const src, \
364
                                const ptrdiff_t src_stride, \
365
                                const int w, const int h, \
366
                                const int mx, const int my \
367
580k
                                HIGHBD_DECL_SUFFIX) \
368
580k
{ \
369
580k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
580k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
580k
} \
mc_tmpl.c:put_8tap_regular_c
Line
Count
Source
367
183k
                                HIGHBD_DECL_SUFFIX) \
368
183k
{ \
369
183k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
183k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
183k
} \
mc_tmpl.c:put_8tap_regular_smooth_c
Line
Count
Source
367
5.61k
                                HIGHBD_DECL_SUFFIX) \
368
5.61k
{ \
369
5.61k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
5.61k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
5.61k
} \
mc_tmpl.c:put_8tap_regular_sharp_c
Line
Count
Source
367
1.57k
                                HIGHBD_DECL_SUFFIX) \
368
1.57k
{ \
369
1.57k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
1.57k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
1.57k
} \
mc_tmpl.c:put_8tap_sharp_regular_c
Line
Count
Source
367
1.88k
                                HIGHBD_DECL_SUFFIX) \
368
1.88k
{ \
369
1.88k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
1.88k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
1.88k
} \
mc_tmpl.c:put_8tap_sharp_smooth_c
Line
Count
Source
367
1.03k
                                HIGHBD_DECL_SUFFIX) \
368
1.03k
{ \
369
1.03k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
1.03k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
1.03k
} \
mc_tmpl.c:put_8tap_sharp_c
Line
Count
Source
367
36.1k
                                HIGHBD_DECL_SUFFIX) \
368
36.1k
{ \
369
36.1k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
36.1k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
36.1k
} \
mc_tmpl.c:put_8tap_smooth_regular_c
Line
Count
Source
367
4.37k
                                HIGHBD_DECL_SUFFIX) \
368
4.37k
{ \
369
4.37k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
4.37k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
4.37k
} \
mc_tmpl.c:put_8tap_smooth_c
Line
Count
Source
367
345k
                                HIGHBD_DECL_SUFFIX) \
368
345k
{ \
369
345k
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
345k
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
345k
} \
mc_tmpl.c:put_8tap_smooth_sharp_c
Line
Count
Source
367
836
                                HIGHBD_DECL_SUFFIX) \
368
836
{ \
369
836
    put_8tap_c(dst, dst_stride, src, src_stride, w, h, mx, my, \
370
836
               type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
371
836
} \
372
static void put_8tap_##type##_scaled_c(pixel *const dst, \
373
                                       const ptrdiff_t dst_stride, \
374
                                       const pixel *const src, \
375
                                       const ptrdiff_t src_stride, \
376
                                       const int w, const int h, \
377
                                       const int mx, const int my, \
378
                                       const int dx, const int dy \
379
268k
                                       HIGHBD_DECL_SUFFIX) \
380
268k
{ \
381
268k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
268k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
268k
} \
mc_tmpl.c:put_8tap_regular_scaled_c
Line
Count
Source
379
189k
                                       HIGHBD_DECL_SUFFIX) \
380
189k
{ \
381
189k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
189k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
189k
} \
mc_tmpl.c:put_8tap_regular_smooth_scaled_c
Line
Count
Source
379
11.9k
                                       HIGHBD_DECL_SUFFIX) \
380
11.9k
{ \
381
11.9k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
11.9k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
11.9k
} \
mc_tmpl.c:put_8tap_regular_sharp_scaled_c
Line
Count
Source
379
1.39k
                                       HIGHBD_DECL_SUFFIX) \
380
1.39k
{ \
381
1.39k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
1.39k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
1.39k
} \
mc_tmpl.c:put_8tap_sharp_regular_scaled_c
Line
Count
Source
379
2.91k
                                       HIGHBD_DECL_SUFFIX) \
380
2.91k
{ \
381
2.91k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
2.91k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
2.91k
} \
mc_tmpl.c:put_8tap_sharp_smooth_scaled_c
Line
Count
Source
379
1.05k
                                       HIGHBD_DECL_SUFFIX) \
380
1.05k
{ \
381
1.05k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
1.05k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
1.05k
} \
mc_tmpl.c:put_8tap_sharp_scaled_c
Line
Count
Source
379
21.9k
                                       HIGHBD_DECL_SUFFIX) \
380
21.9k
{ \
381
21.9k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
21.9k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
21.9k
} \
mc_tmpl.c:put_8tap_smooth_regular_scaled_c
Line
Count
Source
379
16.1k
                                       HIGHBD_DECL_SUFFIX) \
380
16.1k
{ \
381
16.1k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
16.1k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
16.1k
} \
mc_tmpl.c:put_8tap_smooth_scaled_c
Line
Count
Source
379
22.6k
                                       HIGHBD_DECL_SUFFIX) \
380
22.6k
{ \
381
22.6k
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
22.6k
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
22.6k
} \
mc_tmpl.c:put_8tap_smooth_sharp_scaled_c
Line
Count
Source
379
977
                                       HIGHBD_DECL_SUFFIX) \
380
977
{ \
381
977
    put_8tap_scaled_c(dst, dst_stride, src, src_stride, w, h, mx, my, dx, dy, \
382
977
                      type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
383
977
} \
384
static void prep_8tap_##type##_c(int16_t *const tmp, \
385
                                 const pixel *const src, \
386
                                 const ptrdiff_t src_stride, \
387
                                 const int w, const int h, \
388
                                 const int mx, const int my \
389
136k
                                 HIGHBD_DECL_SUFFIX) \
390
136k
{ \
391
136k
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
136k
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
136k
} \
mc_tmpl.c:prep_8tap_regular_c
Line
Count
Source
389
31.9k
                                 HIGHBD_DECL_SUFFIX) \
390
31.9k
{ \
391
31.9k
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
31.9k
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
31.9k
} \
mc_tmpl.c:prep_8tap_regular_smooth_c
Line
Count
Source
389
963
                                 HIGHBD_DECL_SUFFIX) \
390
963
{ \
391
963
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
963
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
963
} \
mc_tmpl.c:prep_8tap_regular_sharp_c
Line
Count
Source
389
2.99k
                                 HIGHBD_DECL_SUFFIX) \
390
2.99k
{ \
391
2.99k
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
2.99k
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
2.99k
} \
mc_tmpl.c:prep_8tap_sharp_regular_c
Line
Count
Source
389
5.87k
                                 HIGHBD_DECL_SUFFIX) \
390
5.87k
{ \
391
5.87k
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
5.87k
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
5.87k
} \
mc_tmpl.c:prep_8tap_sharp_smooth_c
Line
Count
Source
389
1.56k
                                 HIGHBD_DECL_SUFFIX) \
390
1.56k
{ \
391
1.56k
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
1.56k
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
1.56k
} \
mc_tmpl.c:prep_8tap_sharp_c
Line
Count
Source
389
23.7k
                                 HIGHBD_DECL_SUFFIX) \
390
23.7k
{ \
391
23.7k
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
23.7k
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
23.7k
} \
mc_tmpl.c:prep_8tap_smooth_regular_c
Line
Count
Source
389
990
                                 HIGHBD_DECL_SUFFIX) \
390
990
{ \
391
990
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
990
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
990
} \
mc_tmpl.c:prep_8tap_smooth_c
Line
Count
Source
389
67.6k
                                 HIGHBD_DECL_SUFFIX) \
390
67.6k
{ \
391
67.6k
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
67.6k
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
67.6k
} \
mc_tmpl.c:prep_8tap_smooth_sharp_c
Line
Count
Source
389
1.02k
                                 HIGHBD_DECL_SUFFIX) \
390
1.02k
{ \
391
1.02k
    prep_8tap_c(tmp, src, src_stride, w, h, mx, my, \
392
1.02k
                type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
393
1.02k
} \
394
static void prep_8tap_##type##_scaled_c(int16_t *const tmp, \
395
                                        const pixel *const src, \
396
                                        const ptrdiff_t src_stride, \
397
                                        const int w, const int h, \
398
                                        const int mx, const int my, \
399
                                        const int dx, const int dy \
400
88.0k
                                        HIGHBD_DECL_SUFFIX) \
401
88.0k
{ \
402
88.0k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
88.0k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
88.0k
}
mc_tmpl.c:prep_8tap_regular_scaled_c
Line
Count
Source
400
25.1k
                                        HIGHBD_DECL_SUFFIX) \
401
25.1k
{ \
402
25.1k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
25.1k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
25.1k
}
mc_tmpl.c:prep_8tap_regular_smooth_scaled_c
Line
Count
Source
400
2.46k
                                        HIGHBD_DECL_SUFFIX) \
401
2.46k
{ \
402
2.46k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
2.46k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
2.46k
}
mc_tmpl.c:prep_8tap_regular_sharp_scaled_c
Line
Count
Source
400
2.60k
                                        HIGHBD_DECL_SUFFIX) \
401
2.60k
{ \
402
2.60k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
2.60k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
2.60k
}
mc_tmpl.c:prep_8tap_sharp_regular_scaled_c
Line
Count
Source
400
12.9k
                                        HIGHBD_DECL_SUFFIX) \
401
12.9k
{ \
402
12.9k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
12.9k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
12.9k
}
mc_tmpl.c:prep_8tap_sharp_smooth_scaled_c
Line
Count
Source
400
1.65k
                                        HIGHBD_DECL_SUFFIX) \
401
1.65k
{ \
402
1.65k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
1.65k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
1.65k
}
mc_tmpl.c:prep_8tap_sharp_scaled_c
Line
Count
Source
400
5.86k
                                        HIGHBD_DECL_SUFFIX) \
401
5.86k
{ \
402
5.86k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
5.86k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
5.86k
}
mc_tmpl.c:prep_8tap_smooth_regular_scaled_c
Line
Count
Source
400
5.75k
                                        HIGHBD_DECL_SUFFIX) \
401
5.75k
{ \
402
5.75k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
5.75k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
5.75k
}
mc_tmpl.c:prep_8tap_smooth_scaled_c
Line
Count
Source
400
29.4k
                                        HIGHBD_DECL_SUFFIX) \
401
29.4k
{ \
402
29.4k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
29.4k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
29.4k
}
mc_tmpl.c:prep_8tap_smooth_sharp_scaled_c
Line
Count
Source
400
2.09k
                                        HIGHBD_DECL_SUFFIX) \
401
2.09k
{ \
402
2.09k
    prep_8tap_scaled_c(tmp, src, src_stride, w, h, mx, my, dx, dy, \
403
2.09k
                       type_h | (type_v << 2) HIGHBD_TAIL_SUFFIX); \
404
2.09k
}
405
406
filter_fns(regular,        DAV1D_FILTER_8TAP_REGULAR, DAV1D_FILTER_8TAP_REGULAR)
407
filter_fns(regular_sharp,  DAV1D_FILTER_8TAP_REGULAR, DAV1D_FILTER_8TAP_SHARP)
408
filter_fns(regular_smooth, DAV1D_FILTER_8TAP_REGULAR, DAV1D_FILTER_8TAP_SMOOTH)
409
filter_fns(smooth,         DAV1D_FILTER_8TAP_SMOOTH,  DAV1D_FILTER_8TAP_SMOOTH)
410
filter_fns(smooth_regular, DAV1D_FILTER_8TAP_SMOOTH,  DAV1D_FILTER_8TAP_REGULAR)
411
filter_fns(smooth_sharp,   DAV1D_FILTER_8TAP_SMOOTH,  DAV1D_FILTER_8TAP_SHARP)
412
filter_fns(sharp,          DAV1D_FILTER_8TAP_SHARP,   DAV1D_FILTER_8TAP_SHARP)
413
filter_fns(sharp_regular,  DAV1D_FILTER_8TAP_SHARP,   DAV1D_FILTER_8TAP_REGULAR)
414
filter_fns(sharp_smooth,   DAV1D_FILTER_8TAP_SHARP,   DAV1D_FILTER_8TAP_SMOOTH)
415
416
#define FILTER_BILIN(src, x, mxy, stride) \
417
13.8M
    (16 * src[x] + ((mxy) * (src[x + stride] - src[x])))
418
419
#define FILTER_BILIN_RND(src, x, mxy, stride, sh) \
420
13.8M
    ((FILTER_BILIN(src, x, mxy, stride) + ((1 << (sh)) >> 1)) >> (sh))
421
422
#define FILTER_BILIN_CLIP(src, x, mxy, stride, sh) \
423
1.53M
    iclip_pixel(FILTER_BILIN_RND(src, x, mxy, stride, sh))
424
425
#define FILTER_BILIN2(src1, src2, x, mxy) \
426
5.36M
    (16 * src1[x] + ((mxy) * (src2[x] - src1[x])))
427
428
#define FILTER_BILIN_RND2(src1, src2, x, mxy, sh) \
429
5.36M
    ((FILTER_BILIN2(src1, src2, x, mxy) + ((1 << (sh)) >> 1)) >> (sh))
430
431
#define FILTER_BILIN_CLIP2(src1, src2, x, mxy, sh) \
432
4.21M
    iclip_pixel(FILTER_BILIN_RND2(src1, src2, x, mxy, sh))
433
434
static void put_bilin_c(pixel *dst, ptrdiff_t dst_stride,
435
                        const pixel *src, ptrdiff_t src_stride,
436
                        const int w, int h, const int mx, const int my
437
                        HIGHBD_DECL_SUFFIX)
438
144k
{
439
144k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
440
144k
    const int intermediate_rnd = (1 << intermediate_bits) >> 1;
441
144k
    dst_stride = PXSTRIDE(dst_stride);
442
144k
    src_stride = PXSTRIDE(src_stride);
443
444
144k
    if (mx) {
445
16.1k
        if (my) {
446
10.7k
            int16_t mid[128 * 129], *mid_ptr = mid;
447
10.7k
            int tmp_h = h + 1;
448
449
89.0k
            do {
450
1.18M
                for (int x = 0; x < w; x++)
451
1.09M
                    mid_ptr[x] = FILTER_BILIN_RND(src, x, mx, 1,
452
89.0k
                                                  4 - intermediate_bits);
453
454
89.0k
                mid_ptr += 128;
455
89.0k
                src += src_stride;
456
89.0k
            } while (--tmp_h);
457
458
10.7k
            mid_ptr = mid;
459
78.3k
            do {
460
1.10M
                for (int x = 0; x < w; x++)
461
1.02M
                    dst[x] = FILTER_BILIN_CLIP(mid_ptr, x, my, 128,
462
78.3k
                                               4 + intermediate_bits);
463
464
78.3k
                mid_ptr += 128;
465
78.3k
                dst += dst_stride;
466
78.3k
            } while (--h);
467
10.7k
        } else {
468
44.4k
            do {
469
583k
                for (int x = 0; x < w; x++) {
470
538k
                    const int px = FILTER_BILIN_RND(src, x, mx, 1,
471
538k
                                                    4 - intermediate_bits);
472
538k
                    dst[x] = iclip_pixel((px + intermediate_rnd) >> intermediate_bits);
473
538k
                }
474
475
44.4k
                dst += dst_stride;
476
44.4k
                src += src_stride;
477
44.4k
            } while (--h);
478
5.34k
        }
479
128k
    } else if (my) {
480
28.7k
        do {
481
537k
            for (int x = 0; x < w; x++)
482
508k
                dst[x] = FILTER_BILIN_CLIP(src, x, my, src_stride, 4);
483
484
28.7k
            dst += dst_stride;
485
28.7k
            src += src_stride;
486
28.7k
        } while (--h);
487
3.13k
    } else
488
124k
        put_c(dst, dst_stride, src, src_stride, w, h);
489
144k
}
490
491
static void put_bilin_scaled_c(pixel *dst, ptrdiff_t dst_stride,
492
                               const pixel *src, ptrdiff_t src_stride,
493
                               const int w, int h, const int mx, int my,
494
                               const int dx, const int dy
495
                               HIGHBD_DECL_SUFFIX)
496
14.3k
{
497
14.3k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
498
14.3k
    int16_t mid[128 * 2];
499
14.3k
    int in_y = -2;
500
501
202k
    do {
502
202k
        int x;
503
202k
        int y = my >> 10;
504
202k
        int16_t *mid1 = &mid[(y & 1) * 128];
505
202k
        int16_t *mid2 = &mid[((y + 1) & 1) * 128];
506
202k
        int dmy = my & 0x3ff;
507
508
424k
        while (in_y < y) {
509
221k
            int imx = mx, ioff = 0;
510
221k
            int16_t *mid_ptr = &mid[(in_y & 1) * 128];
511
512
4.79M
            for (x = 0; x < w; x++) {
513
4.57M
                mid_ptr[x] = FILTER_BILIN_RND(src, ioff, imx >> 6, 1,
514
4.57M
                                              4 - intermediate_bits);
515
4.57M
                imx += dx;
516
4.57M
                ioff += imx >> 10;
517
4.57M
                imx &= 0x3ff;
518
4.57M
            }
519
520
221k
            src += PXSTRIDE(src_stride);
521
221k
            in_y++;
522
221k
        }
523
524
4.42M
        for (x = 0; x < w; x++)
525
4.21M
            dst[x] = FILTER_BILIN_CLIP2(mid1, mid2, x, dmy >> 6,
526
202k
                                       4 + intermediate_bits);
527
528
202k
        my += dy;
529
202k
        dst += PXSTRIDE(dst_stride);
530
202k
    } while (--h);
531
14.3k
}
532
533
static void prep_bilin_c(int16_t *tmp,
534
                         const pixel *src, ptrdiff_t src_stride,
535
                         const int w, int h, const int mx, const int my
536
                         HIGHBD_DECL_SUFFIX)
537
13.0k
{
538
13.0k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
539
13.0k
    src_stride = PXSTRIDE(src_stride);
540
541
13.0k
    if (mx) {
542
6.52k
        if (my) {
543
4.60k
            int16_t mid[128 * 129], *mid_ptr = mid;
544
4.60k
            int tmp_h = h + 1;
545
546
86.2k
            do {
547
1.64M
                for (int x = 0; x < w; x++)
548
1.56M
                    mid_ptr[x] = FILTER_BILIN_RND(src, x, mx, 1,
549
86.2k
                                                  4 - intermediate_bits);
550
551
86.2k
                mid_ptr += 128;
552
86.2k
                src += src_stride;
553
86.2k
            } while (--tmp_h);
554
555
4.60k
            mid_ptr = mid;
556
84.7k
            do {
557
1.66M
                for (int x = 0; x < w; x++)
558
1.57M
                    tmp[x] = FILTER_BILIN_RND(mid_ptr, x, my, 128, 4) -
559
1.57M
                             PREP_BIAS;
560
561
84.7k
                mid_ptr += 128;
562
84.7k
                tmp += w;
563
84.7k
            } while (--h);
564
4.60k
        } else {
565
39.9k
            do {
566
975k
                for (int x = 0; x < w; x++)
567
935k
                    tmp[x] = FILTER_BILIN_RND(src, x, mx, 1,
568
935k
                                              4 - intermediate_bits) -
569
935k
                             PREP_BIAS;
570
571
39.9k
                tmp += w;
572
39.9k
                src += src_stride;
573
39.9k
            } while (--h);
574
1.92k
        }
575
6.52k
    } else if (my) {
576
48.4k
        do {
577
1.00M
            for (int x = 0; x < w; x++)
578
957k
                tmp[x] = FILTER_BILIN_RND(src, x, my, src_stride,
579
957k
                                          4 - intermediate_bits) - PREP_BIAS;
580
581
48.4k
            tmp += w;
582
48.4k
            src += src_stride;
583
48.4k
        } while (--h);
584
2.19k
    } else
585
4.30k
        prep_c(tmp, src, src_stride, w, h HIGHBD_TAIL_SUFFIX);
586
13.0k
}
587
588
static void prep_bilin_scaled_c(int16_t *tmp,
589
                                const pixel *src, ptrdiff_t src_stride,
590
                                const int w, int h, const int mx, int my,
591
                                const int dx, const int dy HIGHBD_DECL_SUFFIX)
592
2.52k
{
593
2.52k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
594
2.52k
    int16_t mid[128 * 2];
595
2.52k
    int in_y = -2;
596
597
48.6k
    do {
598
48.6k
        int x;
599
48.6k
        int y = my >> 10;
600
48.6k
        int16_t *mid1 = &mid[(y & 1) * 128];
601
48.6k
        int16_t *mid2 = &mid[((y + 1) & 1) * 128];
602
48.6k
        int dmy = my & 0x3ff;
603
604
95.3k
        while (in_y < y) {
605
46.7k
            int imx = mx, ioff = 0;
606
46.7k
            int16_t *mid_ptr = &mid[(in_y & 1) * 128];
607
608
1.16M
            for (x = 0; x < w; x++) {
609
1.11M
                mid_ptr[x] = FILTER_BILIN_RND(src, ioff, imx >> 6, 1,
610
1.11M
                                              4 - intermediate_bits);
611
1.11M
                imx += dx;
612
1.11M
                ioff += imx >> 10;
613
1.11M
                imx &= 0x3ff;
614
1.11M
            }
615
616
46.7k
            src += PXSTRIDE(src_stride);
617
46.7k
            in_y++;
618
46.7k
        }
619
620
1.19M
        for (x = 0; x < w; x++)
621
1.14M
            tmp[x] = FILTER_BILIN_RND2(mid1, mid2, x, dmy >> 6, 4) - PREP_BIAS;
622
623
48.6k
        my += dy;
624
48.6k
        tmp += w;
625
48.6k
    } while (--h);
626
2.52k
}
627
628
static void avg_c(pixel *dst, const ptrdiff_t dst_stride,
629
                  const int16_t *tmp1, const int16_t *tmp2, const int w, int h
630
                  HIGHBD_DECL_SUFFIX)
631
84.3k
{
632
84.3k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
633
84.3k
    const int sh = intermediate_bits + 1;
634
84.3k
    const int rnd = (1 << intermediate_bits) + PREP_BIAS * 2;
635
1.82M
    do {
636
47.0M
        for (int x = 0; x < w; x++)
637
45.2M
            dst[x] = iclip_pixel((tmp1[x] + tmp2[x] + rnd) >> sh);
638
639
1.82M
        tmp1 += w;
640
1.82M
        tmp2 += w;
641
1.82M
        dst += PXSTRIDE(dst_stride);
642
1.82M
    } while (--h);
643
84.3k
}
644
645
static void w_avg_c(pixel *dst, const ptrdiff_t dst_stride,
646
                    const int16_t *tmp1, const int16_t *tmp2, const int w, int h,
647
                    const int weight HIGHBD_DECL_SUFFIX)
648
21.7k
{
649
21.7k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
650
21.7k
    const int sh = intermediate_bits + 4;
651
21.7k
    const int rnd = (8 << intermediate_bits) + PREP_BIAS * 16;
652
261k
    do {
653
4.91M
        for (int x = 0; x < w; x++)
654
4.65M
            dst[x] = iclip_pixel((tmp1[x] * weight +
655
4.65M
                                  tmp2[x] * (16 - weight) + rnd) >> sh);
656
657
261k
        tmp1 += w;
658
261k
        tmp2 += w;
659
261k
        dst += PXSTRIDE(dst_stride);
660
261k
    } while (--h);
661
21.7k
}
662
663
static void mask_c(pixel *dst, const ptrdiff_t dst_stride,
664
                   const int16_t *tmp1, const int16_t *tmp2, const int w, int h,
665
                   const uint8_t *mask HIGHBD_DECL_SUFFIX)
666
14.5k
{
667
14.5k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
668
14.5k
    const int sh = intermediate_bits + 6;
669
14.5k
    const int rnd = (32 << intermediate_bits) + PREP_BIAS * 64;
670
238k
    do {
671
5.05M
        for (int x = 0; x < w; x++)
672
4.81M
            dst[x] = iclip_pixel((tmp1[x] * mask[x] +
673
4.81M
                                  tmp2[x] * (64 - mask[x]) + rnd) >> sh);
674
675
238k
        tmp1 += w;
676
238k
        tmp2 += w;
677
238k
        mask += w;
678
238k
        dst += PXSTRIDE(dst_stride);
679
238k
    } while (--h);
680
14.5k
}
681
682
13.7M
#define blend_px(a, b, m) (((a * (64 - m) + b * m) + 32) >> 6)
683
static void blend_c(pixel *dst, const ptrdiff_t dst_stride, const pixel *tmp,
684
                    const int w, int h, const uint8_t *mask)
685
20.3k
{
686
205k
    do {
687
2.76M
        for (int x = 0; x < w; x++) {
688
2.55M
            dst[x] = blend_px(dst[x], tmp[x], mask[x]);
689
2.55M
        }
690
205k
        dst += PXSTRIDE(dst_stride);
691
205k
        tmp += w;
692
205k
        mask += w;
693
205k
    } while (--h);
694
20.3k
}
695
696
static void blend_v_c(pixel *dst, const ptrdiff_t dst_stride, const pixel *tmp,
697
                      const int w, int h)
698
65.5k
{
699
65.5k
    const uint8_t *const mask = &dav1d_obmc_masks[w];
700
787k
    do {
701
5.21M
        for (int x = 0; x < (w * 3) >> 2; x++) {
702
4.42M
            dst[x] = blend_px(dst[x], tmp[x], mask[x]);
703
4.42M
        }
704
787k
        dst += PXSTRIDE(dst_stride);
705
787k
        tmp += w;
706
787k
    } while (--h);
707
65.5k
}
708
709
static void blend_h_c(pixel *dst, const ptrdiff_t dst_stride, const pixel *tmp,
710
                      const int w, int h)
711
94.3k
{
712
94.3k
    const uint8_t *mask = &dav1d_obmc_masks[h];
713
94.3k
    h = (h * 3) >> 2;
714
565k
    do {
715
565k
        const int m = *mask++;
716
7.34M
        for (int x = 0; x < w; x++) {
717
6.78M
            dst[x] = blend_px(dst[x], tmp[x], m);
718
6.78M
        }
719
565k
        dst += PXSTRIDE(dst_stride);
720
565k
        tmp += w;
721
565k
    } while (--h);
722
94.3k
}
723
724
static void w_mask_c(pixel *dst, const ptrdiff_t dst_stride,
725
                     const int16_t *tmp1, const int16_t *tmp2, const int w, int h,
726
                     uint8_t *mask, const int sign,
727
                     const int ss_hor, const int ss_ver HIGHBD_DECL_SUFFIX)
728
5.05k
{
729
    // store mask at 2x2 resolution, i.e. store 2x1 sum for even rows,
730
    // and then load this intermediate to calculate final value for odd rows
731
5.05k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
732
5.05k
    const int bitdepth = bitdepth_from_max(bitdepth_max);
733
5.05k
    const int sh = intermediate_bits + 6;
734
5.05k
    const int rnd = (32 << intermediate_bits) + PREP_BIAS * 64;
735
5.05k
    const int mask_sh = bitdepth + intermediate_bits - 4;
736
5.05k
    const int mask_rnd = 1 << (mask_sh - 5);
737
160k
    do {
738
4.35M
        for (int x = 0; x < w; x++) {
739
4.19M
            const int tmpdiff = tmp1[x] - tmp2[x];
740
4.19M
            const int m = imin(38 + ((abs(tmpdiff) + mask_rnd) >> mask_sh), 64);
741
4.19M
            dst[x] = iclip_pixel((tmpdiff * m + tmp2[x] * 64 + rnd) >> sh);
742
743
4.19M
            if (ss_hor) {
744
2.41M
                x++;
745
746
2.41M
                const int tmpdiff = tmp1[x] - tmp2[x];
747
2.41M
                const int n = imin(38 + ((abs(tmpdiff) + mask_rnd) >> mask_sh), 64);
748
2.41M
                dst[x] = iclip_pixel((tmpdiff * n + tmp2[x] * 64 + rnd) >> sh);
749
750
2.41M
                if (h & ss_ver) {
751
854k
                    mask[x >> 1] = (m + n + mask[x >> 1] + 2 - sign) >> 2;
752
1.56M
                } else if (ss_ver) {
753
854k
                    mask[x >> 1] = m + n;
754
854k
                } else {
755
706k
                    mask[x >> 1] = (m + n + 1 - sign) >> 1;
756
706k
                }
757
2.41M
            } else {
758
1.77M
                mask[x] = m;
759
1.77M
            }
760
4.19M
        }
761
762
160k
        tmp1 += w;
763
160k
        tmp2 += w;
764
160k
        dst += PXSTRIDE(dst_stride);
765
160k
        if (!ss_ver || (h & 1)) mask += w >> ss_hor;
766
160k
    } while (--h);
767
5.05k
}
768
769
#define w_mask_fns(ssn, ss_hor, ss_ver) \
770
static void w_mask_##ssn##_c(pixel *const dst, const ptrdiff_t dst_stride, \
771
                             const int16_t *const tmp1, const int16_t *const tmp2, \
772
                             const int w, const int h, uint8_t *mask, \
773
5.06k
                             const int sign HIGHBD_DECL_SUFFIX) \
774
5.06k
{ \
775
5.06k
    w_mask_c(dst, dst_stride, tmp1, tmp2, w, h, mask, sign, ss_hor, ss_ver \
776
5.06k
             HIGHBD_TAIL_SUFFIX); \
777
5.06k
}
mc_tmpl.c:w_mask_444_c
Line
Count
Source
773
1.03k
                             const int sign HIGHBD_DECL_SUFFIX) \
774
1.03k
{ \
775
1.03k
    w_mask_c(dst, dst_stride, tmp1, tmp2, w, h, mask, sign, ss_hor, ss_ver \
776
1.03k
             HIGHBD_TAIL_SUFFIX); \
777
1.03k
}
mc_tmpl.c:w_mask_422_c
Line
Count
Source
773
1.13k
                             const int sign HIGHBD_DECL_SUFFIX) \
774
1.13k
{ \
775
1.13k
    w_mask_c(dst, dst_stride, tmp1, tmp2, w, h, mask, sign, ss_hor, ss_ver \
776
1.13k
             HIGHBD_TAIL_SUFFIX); \
777
1.13k
}
mc_tmpl.c:w_mask_420_c
Line
Count
Source
773
2.89k
                             const int sign HIGHBD_DECL_SUFFIX) \
774
2.89k
{ \
775
2.89k
    w_mask_c(dst, dst_stride, tmp1, tmp2, w, h, mask, sign, ss_hor, ss_ver \
776
2.89k
             HIGHBD_TAIL_SUFFIX); \
777
2.89k
}
778
779
w_mask_fns(444, 0, 0);
780
w_mask_fns(422, 1, 0);
781
w_mask_fns(420, 1, 1);
782
783
#undef w_mask_fns
784
785
#define FILTER_WARP_RND(src, x, F, stride, sh) \
786
60.8M
    ((F[0] * src[x - 3 * stride] + \
787
60.8M
      F[1] * src[x - 2 * stride] + \
788
60.8M
      F[2] * src[x - 1 * stride] + \
789
60.8M
      F[3] * src[x + 0 * stride] + \
790
60.8M
      F[4] * src[x + 1 * stride] + \
791
60.8M
      F[5] * src[x + 2 * stride] + \
792
60.8M
      F[6] * src[x + 3 * stride] + \
793
60.8M
      F[7] * src[x + 4 * stride] + \
794
60.8M
      ((1 << (sh)) >> 1)) >> (sh))
795
796
#define FILTER_WARP_CLIP(src, x, F, stride, sh) \
797
12.8M
    iclip_pixel(FILTER_WARP_RND(src, x, F, stride, sh))
798
799
static void warp_affine_8x8_c(pixel *dst, const ptrdiff_t dst_stride,
800
                              const pixel *src, const ptrdiff_t src_stride,
801
                              const int16_t *const abcd, int mx, int my
802
                              HIGHBD_DECL_SUFFIX)
803
217k
{
804
217k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
805
217k
    int16_t mid[15 * 8], *mid_ptr = mid;
806
807
217k
    src -= 3 * PXSTRIDE(src_stride);
808
3.41M
    for (int y = 0; y < 15; y++, mx += abcd[1]) {
809
28.4M
        for (int x = 0, tmx = mx; x < 8; x++, tmx += abcd[0]) {
810
25.2M
            const int8_t *const filter =
811
25.2M
                dav1d_mc_warp_filter[64 + ((tmx + 512) >> 10)];
812
813
25.2M
            mid_ptr[x] = FILTER_WARP_RND(src, x, filter, 1,
814
25.2M
                                         7 - intermediate_bits);
815
25.2M
        }
816
3.19M
        src += PXSTRIDE(src_stride);
817
3.19M
        mid_ptr += 8;
818
3.19M
    }
819
820
217k
    mid_ptr = &mid[3 * 8];
821
1.83M
    for (int y = 0; y < 8; y++, my += abcd[3]) {
822
14.4M
        for (int x = 0, tmy = my; x < 8; x++, tmy += abcd[2]) {
823
12.8M
            const int8_t *const filter =
824
12.8M
                dav1d_mc_warp_filter[64 + ((tmy + 512) >> 10)];
825
826
12.8M
            dst[x] = FILTER_WARP_CLIP(mid_ptr, x, filter, 8,
827
12.8M
                                      7 + intermediate_bits);
828
12.8M
        }
829
1.61M
        mid_ptr += 8;
830
1.61M
        dst += PXSTRIDE(dst_stride);
831
1.61M
    }
832
217k
}
833
834
static void warp_affine_8x8t_c(int16_t *tmp, const ptrdiff_t tmp_stride,
835
                               const pixel *src, const ptrdiff_t src_stride,
836
                               const int16_t *const abcd, int mx, int my
837
                               HIGHBD_DECL_SUFFIX)
838
124k
{
839
124k
    const int intermediate_bits = get_intermediate_bits(bitdepth_max);
840
124k
    int16_t mid[15 * 8], *mid_ptr = mid;
841
842
124k
    src -= 3 * PXSTRIDE(src_stride);
843
1.98M
    for (int y = 0; y < 15; y++, mx += abcd[1]) {
844
16.6M
        for (int x = 0, tmx = mx; x < 8; x++, tmx += abcd[0]) {
845
14.8M
            const int8_t *const filter =
846
14.8M
                dav1d_mc_warp_filter[64 + ((tmx + 512) >> 10)];
847
848
14.8M
            mid_ptr[x] = FILTER_WARP_RND(src, x, filter, 1,
849
14.8M
                                         7 - intermediate_bits);
850
14.8M
        }
851
1.85M
        src += PXSTRIDE(src_stride);
852
1.85M
        mid_ptr += 8;
853
1.85M
    }
854
855
124k
    mid_ptr = &mid[3 * 8];
856
1.11M
    for (int y = 0; y < 8; y++, my += abcd[3]) {
857
8.92M
        for (int x = 0, tmy = my; x < 8; x++, tmy += abcd[2]) {
858
7.92M
            const int8_t *const filter =
859
7.92M
                dav1d_mc_warp_filter[64 + ((tmy + 512) >> 10)];
860
861
7.92M
            tmp[x] = FILTER_WARP_RND(mid_ptr, x, filter, 8, 7) - PREP_BIAS;
862
7.92M
        }
863
992k
        mid_ptr += 8;
864
992k
        tmp += tmp_stride;
865
992k
    }
866
124k
}
867
868
static void emu_edge_c(const intptr_t bw, const intptr_t bh,
869
                       const intptr_t iw, const intptr_t ih,
870
                       const intptr_t x, const intptr_t y,
871
                       pixel *dst, const ptrdiff_t dst_stride,
872
                       const pixel *ref, const ptrdiff_t ref_stride)
873
1.29M
{
874
    // find offset in reference of visible block to copy
875
1.29M
    ref += iclip((int) y, 0, (int) ih - 1) * PXSTRIDE(ref_stride) +
876
1.29M
           iclip((int) x, 0, (int) iw - 1);
877
878
    // number of pixels to extend (left, right, top, bottom)
879
1.29M
    const int left_ext = iclip((int) -x, 0, (int) bw - 1);
880
1.29M
    const int right_ext = iclip((int) (x + bw - iw), 0, (int) bw - 1);
881
1.29M
    assert(left_ext + right_ext < bw);
882
1.29M
    const int top_ext = iclip((int) -y, 0, (int) bh - 1);
883
1.29M
    const int bottom_ext = iclip((int) (y + bh - ih), 0, (int) bh - 1);
884
1.29M
    assert(top_ext + bottom_ext < bh);
885
886
    // copy visible portion first
887
1.29M
    pixel *blk = dst + top_ext * PXSTRIDE(dst_stride);
888
1.29M
    const int center_w = (int) (bw - left_ext - right_ext);
889
1.29M
    const int center_h = (int) (bh - top_ext - bottom_ext);
890
19.1M
    for (int y = 0; y < center_h; y++) {
891
17.8M
        pixel_copy(blk + left_ext, ref, center_w);
892
        // extend left edge for this line
893
17.8M
        if (left_ext)
894
4.53M
            pixel_set(blk, blk[left_ext], left_ext);
895
        // extend right edge for this line
896
17.8M
        if (right_ext)
897
14.6M
            pixel_set(blk + left_ext + center_w, blk[left_ext + center_w - 1],
898
14.6M
                      right_ext);
899
17.8M
        ref += PXSTRIDE(ref_stride);
900
17.8M
        blk += PXSTRIDE(dst_stride);
901
17.8M
    }
902
903
    // copy top
904
1.29M
    blk = dst + top_ext * PXSTRIDE(dst_stride);
905
2.54M
    for (int y = 0; y < top_ext; y++) {
906
1.25M
        pixel_copy(dst, blk, bw);
907
1.25M
        dst += PXSTRIDE(dst_stride);
908
1.25M
    }
909
910
    // copy bottom
911
1.29M
    dst += center_h * PXSTRIDE(dst_stride);
912
5.10M
    for (int y = 0; y < bottom_ext; y++) {
913
3.81M
        pixel_copy(dst, &dst[-PXSTRIDE(dst_stride)], bw);
914
3.81M
        dst += PXSTRIDE(dst_stride);
915
3.81M
    }
916
1.29M
}
917
918
static void resize_c(pixel *dst, const ptrdiff_t dst_stride,
919
                     const pixel *src, const ptrdiff_t src_stride,
920
                     const int dst_w, int h, const int src_w,
921
                     const int dx, const int mx0 HIGHBD_DECL_SUFFIX)
922
106k
{
923
3.50M
    do {
924
3.50M
        int mx = mx0, src_x = -1;
925
182M
        for (int x = 0; x < dst_w; x++) {
926
178M
            const int8_t *const F = dav1d_resize_filter[mx >> 8];
927
178M
            dst[x] = iclip_pixel((-(F[0] * src[iclip(src_x - 3, 0, src_w - 1)] +
928
178M
                                    F[1] * src[iclip(src_x - 2, 0, src_w - 1)] +
929
178M
                                    F[2] * src[iclip(src_x - 1, 0, src_w - 1)] +
930
178M
                                    F[3] * src[iclip(src_x + 0, 0, src_w - 1)] +
931
178M
                                    F[4] * src[iclip(src_x + 1, 0, src_w - 1)] +
932
178M
                                    F[5] * src[iclip(src_x + 2, 0, src_w - 1)] +
933
178M
                                    F[6] * src[iclip(src_x + 3, 0, src_w - 1)] +
934
178M
                                    F[7] * src[iclip(src_x + 4, 0, src_w - 1)]) +
935
178M
                                  64) >> 7);
936
178M
            mx += dx;
937
178M
            src_x += mx >> 14;
938
178M
            mx &= 0x3fff;
939
178M
        }
940
941
3.50M
        dst += PXSTRIDE(dst_stride);
942
3.50M
        src += PXSTRIDE(src_stride);
943
3.50M
    } while (--h);
944
106k
}
945
946
#if HAVE_ASM
947
#if ARCH_AARCH64 || ARCH_ARM
948
#include "src/arm/mc.h"
949
#elif ARCH_LOONGARCH64
950
#include "src/loongarch/mc.h"
951
#elif ARCH_PPC64LE
952
#include "src/ppc/mc.h"
953
#elif ARCH_RISCV
954
#include "src/riscv/mc.h"
955
#elif ARCH_X86
956
#include "src/x86/mc.h"
957
#endif
958
#endif
959
960
27.6k
COLD void bitfn(dav1d_mc_dsp_init)(Dav1dMCDSPContext *const c) {
961
276k
#define init_mc_fns(type, name) do { \
962
276k
    c->mc        [type] = put_##name##_c; \
963
276k
    c->mc_scaled [type] = put_##name##_scaled_c; \
964
276k
    c->mct       [type] = prep_##name##_c; \
965
276k
    c->mct_scaled[type] = prep_##name##_scaled_c; \
966
276k
} while (0)
967
968
27.6k
    init_mc_fns(FILTER_2D_8TAP_REGULAR,        8tap_regular);
969
27.6k
    init_mc_fns(FILTER_2D_8TAP_REGULAR_SMOOTH, 8tap_regular_smooth);
970
27.6k
    init_mc_fns(FILTER_2D_8TAP_REGULAR_SHARP,  8tap_regular_sharp);
971
27.6k
    init_mc_fns(FILTER_2D_8TAP_SHARP_REGULAR,  8tap_sharp_regular);
972
27.6k
    init_mc_fns(FILTER_2D_8TAP_SHARP_SMOOTH,   8tap_sharp_smooth);
973
27.6k
    init_mc_fns(FILTER_2D_8TAP_SHARP,          8tap_sharp);
974
27.6k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH_REGULAR, 8tap_smooth_regular);
975
27.6k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH,         8tap_smooth);
976
27.6k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH_SHARP,   8tap_smooth_sharp);
977
27.6k
    init_mc_fns(FILTER_2D_BILINEAR,            bilin);
978
979
27.6k
    c->avg      = avg_c;
980
27.6k
    c->w_avg    = w_avg_c;
981
27.6k
    c->mask     = mask_c;
982
27.6k
    c->blend    = blend_c;
983
27.6k
    c->blend_v  = blend_v_c;
984
27.6k
    c->blend_h  = blend_h_c;
985
27.6k
    c->w_mask[0] = w_mask_444_c;
986
27.6k
    c->w_mask[1] = w_mask_422_c;
987
27.6k
    c->w_mask[2] = w_mask_420_c;
988
27.6k
    c->warp8x8  = warp_affine_8x8_c;
989
27.6k
    c->warp8x8t = warp_affine_8x8t_c;
990
27.6k
    c->emu_edge = emu_edge_c;
991
27.6k
    c->resize   = resize_c;
992
993
#if HAVE_ASM
994
#if ARCH_AARCH64 || ARCH_ARM
995
    mc_dsp_init_arm(c);
996
#elif ARCH_LOONGARCH64
997
    mc_dsp_init_loongarch(c);
998
#elif ARCH_PPC64LE
999
    mc_dsp_init_ppc(c);
1000
#elif ARCH_RISCV
1001
    mc_dsp_init_riscv(c);
1002
#elif ARCH_X86
1003
    mc_dsp_init_x86(c);
1004
#endif
1005
#endif
1006
27.6k
}
dav1d_mc_dsp_init_8bpc
Line
Count
Source
960
12.5k
COLD void bitfn(dav1d_mc_dsp_init)(Dav1dMCDSPContext *const c) {
961
12.5k
#define init_mc_fns(type, name) do { \
962
12.5k
    c->mc        [type] = put_##name##_c; \
963
12.5k
    c->mc_scaled [type] = put_##name##_scaled_c; \
964
12.5k
    c->mct       [type] = prep_##name##_c; \
965
12.5k
    c->mct_scaled[type] = prep_##name##_scaled_c; \
966
12.5k
} while (0)
967
968
12.5k
    init_mc_fns(FILTER_2D_8TAP_REGULAR,        8tap_regular);
969
12.5k
    init_mc_fns(FILTER_2D_8TAP_REGULAR_SMOOTH, 8tap_regular_smooth);
970
12.5k
    init_mc_fns(FILTER_2D_8TAP_REGULAR_SHARP,  8tap_regular_sharp);
971
12.5k
    init_mc_fns(FILTER_2D_8TAP_SHARP_REGULAR,  8tap_sharp_regular);
972
12.5k
    init_mc_fns(FILTER_2D_8TAP_SHARP_SMOOTH,   8tap_sharp_smooth);
973
12.5k
    init_mc_fns(FILTER_2D_8TAP_SHARP,          8tap_sharp);
974
12.5k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH_REGULAR, 8tap_smooth_regular);
975
12.5k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH,         8tap_smooth);
976
12.5k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH_SHARP,   8tap_smooth_sharp);
977
12.5k
    init_mc_fns(FILTER_2D_BILINEAR,            bilin);
978
979
12.5k
    c->avg      = avg_c;
980
12.5k
    c->w_avg    = w_avg_c;
981
12.5k
    c->mask     = mask_c;
982
12.5k
    c->blend    = blend_c;
983
12.5k
    c->blend_v  = blend_v_c;
984
12.5k
    c->blend_h  = blend_h_c;
985
12.5k
    c->w_mask[0] = w_mask_444_c;
986
12.5k
    c->w_mask[1] = w_mask_422_c;
987
12.5k
    c->w_mask[2] = w_mask_420_c;
988
12.5k
    c->warp8x8  = warp_affine_8x8_c;
989
12.5k
    c->warp8x8t = warp_affine_8x8t_c;
990
12.5k
    c->emu_edge = emu_edge_c;
991
12.5k
    c->resize   = resize_c;
992
993
#if HAVE_ASM
994
#if ARCH_AARCH64 || ARCH_ARM
995
    mc_dsp_init_arm(c);
996
#elif ARCH_LOONGARCH64
997
    mc_dsp_init_loongarch(c);
998
#elif ARCH_PPC64LE
999
    mc_dsp_init_ppc(c);
1000
#elif ARCH_RISCV
1001
    mc_dsp_init_riscv(c);
1002
#elif ARCH_X86
1003
    mc_dsp_init_x86(c);
1004
#endif
1005
#endif
1006
12.5k
}
dav1d_mc_dsp_init_16bpc
Line
Count
Source
960
15.1k
COLD void bitfn(dav1d_mc_dsp_init)(Dav1dMCDSPContext *const c) {
961
15.1k
#define init_mc_fns(type, name) do { \
962
15.1k
    c->mc        [type] = put_##name##_c; \
963
15.1k
    c->mc_scaled [type] = put_##name##_scaled_c; \
964
15.1k
    c->mct       [type] = prep_##name##_c; \
965
15.1k
    c->mct_scaled[type] = prep_##name##_scaled_c; \
966
15.1k
} while (0)
967
968
15.1k
    init_mc_fns(FILTER_2D_8TAP_REGULAR,        8tap_regular);
969
15.1k
    init_mc_fns(FILTER_2D_8TAP_REGULAR_SMOOTH, 8tap_regular_smooth);
970
15.1k
    init_mc_fns(FILTER_2D_8TAP_REGULAR_SHARP,  8tap_regular_sharp);
971
15.1k
    init_mc_fns(FILTER_2D_8TAP_SHARP_REGULAR,  8tap_sharp_regular);
972
15.1k
    init_mc_fns(FILTER_2D_8TAP_SHARP_SMOOTH,   8tap_sharp_smooth);
973
15.1k
    init_mc_fns(FILTER_2D_8TAP_SHARP,          8tap_sharp);
974
15.1k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH_REGULAR, 8tap_smooth_regular);
975
15.1k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH,         8tap_smooth);
976
15.1k
    init_mc_fns(FILTER_2D_8TAP_SMOOTH_SHARP,   8tap_smooth_sharp);
977
15.1k
    init_mc_fns(FILTER_2D_BILINEAR,            bilin);
978
979
15.1k
    c->avg      = avg_c;
980
15.1k
    c->w_avg    = w_avg_c;
981
15.1k
    c->mask     = mask_c;
982
15.1k
    c->blend    = blend_c;
983
15.1k
    c->blend_v  = blend_v_c;
984
15.1k
    c->blend_h  = blend_h_c;
985
15.1k
    c->w_mask[0] = w_mask_444_c;
986
15.1k
    c->w_mask[1] = w_mask_422_c;
987
15.1k
    c->w_mask[2] = w_mask_420_c;
988
15.1k
    c->warp8x8  = warp_affine_8x8_c;
989
15.1k
    c->warp8x8t = warp_affine_8x8t_c;
990
15.1k
    c->emu_edge = emu_edge_c;
991
15.1k
    c->resize   = resize_c;
992
993
#if HAVE_ASM
994
#if ARCH_AARCH64 || ARCH_ARM
995
    mc_dsp_init_arm(c);
996
#elif ARCH_LOONGARCH64
997
    mc_dsp_init_loongarch(c);
998
#elif ARCH_PPC64LE
999
    mc_dsp_init_ppc(c);
1000
#elif ARCH_RISCV
1001
    mc_dsp_init_riscv(c);
1002
#elif ARCH_X86
1003
    mc_dsp_init_x86(c);
1004
#endif
1005
#endif
1006
15.1k
}