Coverage Report

Created: 2026-08-17 07:50

next uncovered line (L), next uncovered region (R), next uncovered branch (B)
/src/ffmpeg/libswscale/ops_memcpy.c
Line
Count
Source
1
/**
2
 * Copyright (C) 2025 Niklas Haas
3
 *
4
 * This file is part of FFmpeg.
5
 *
6
 * FFmpeg is free software; you can redistribute it and/or
7
 * modify it under the terms of the GNU Lesser General Public
8
 * License as published by the Free Software Foundation; either
9
 * version 2.1 of the License, or (at your option) any later version.
10
 *
11
 * FFmpeg is distributed in the hope that it will be useful,
12
 * but WITHOUT ANY WARRANTY; without even the implied warranty of
13
 * MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE.  See the GNU
14
 * Lesser General Public License for more details.
15
 *
16
 * You should have received a copy of the GNU Lesser General Public
17
 * License along with FFmpeg; if not, write to the Free Software
18
 * Foundation, Inc., 51 Franklin Street, Fifth Floor, Boston, MA 02110-1301 USA
19
 */
20
21
#include "libavutil/avassert.h"
22
#include "libavutil/mem.h"
23
24
#include "ops_internal.h"
25
26
typedef struct MemcpyPriv {
27
    int num_planes;
28
    int index[4]; /* or -1 to clear plane */
29
    uint8_t clear_value[4];
30
} MemcpyPriv;
31
32
/**
33
 * Switch to loop if total padding exceeds this number of bytes. Chosen to
34
 * align with the typical L1 cache size of modern CPUs, as this avoids the
35
 * risk of the implementation loading one extra unnecessary cache line.
36
 */
37
0
#define SWS_MAX_PADDING 64
38
39
/* Memcpy backend for trivial cases */
40
41
static void process(const SwsOpExec *exec, const void *priv,
42
                    int x_start, int y_start, int x_end, int y_end)
43
0
{
44
0
    const MemcpyPriv *p = priv;
45
0
    const int lines = y_end - y_start;
46
0
    av_assert1(x_start == 0 && x_end == exec->width);
47
48
0
    for (int i = 0; i < p->num_planes; i++) {
49
0
        uint8_t *out = exec->out[i];
50
0
        const int idx = p->index[i];
51
0
        const int bytes = x_end * exec->block_size_out[i];
52
0
        const int use_loop = exec->out_stride[i] > bytes + SWS_MAX_PADDING;
53
0
        if (idx < 0 && !use_loop) {
54
0
            memset(out, p->clear_value[i], exec->out_stride[i] * lines);
55
0
        } else if (idx < 0) {
56
0
            for (int y = y_start; y < y_end; y++) {
57
0
                memset(out, p->clear_value[i], bytes);
58
0
                out += exec->out_stride[i];
59
0
            }
60
0
        } else if (out == exec->in[idx]) {
61
0
            av_assert1(exec->out_stride[i] == exec->in_stride[idx]);
62
0
            continue; /* plane was already ref'd */
63
0
        } else if (exec->out_stride[i] == exec->in_stride[idx] && !use_loop) {
64
0
            memcpy(out, exec->in[idx], exec->out_stride[i] * lines);
65
0
        } else {
66
0
            const uint8_t *in = exec->in[idx];
67
0
            for (int y = y_start; y < y_end; y++) {
68
0
                memcpy(out, in, bytes);
69
0
                out += exec->out_stride[i];
70
0
                in  += exec->in_stride[idx];
71
0
            }
72
0
        }
73
0
    }
74
0
}
75
76
static int compile(SwsContext *ctx, const SwsOpList *ops, SwsCompiledOp *out)
77
0
{
78
0
    MemcpyPriv p = {0};
79
80
0
    for (int n = 0; n < ops->num_ops; n++) {
81
0
        const SwsOp *op = &ops->ops[n];
82
0
        switch (op->op) {
83
0
        case SWS_OP_READ:
84
0
            if (ff_sws_rw_op_planes(op) != op->rw.elems || op->rw.frac || op->rw.filter.op)
85
0
                return AVERROR(ENOTSUP);
86
0
            for (int i = 0; i < op->rw.elems; i++)
87
0
                p.index[i] = i;
88
0
            break;
89
90
0
        case SWS_OP_SWIZZLE: {
91
0
            const MemcpyPriv orig = p;
92
0
            for (int i = 0; i < 4; i++) {
93
                /* Explicitly exclude swizzle masks that contain duplicates,
94
                 * because these are wasteful to implement as a memcpy */
95
0
                for (int j = 0; j < i; j++) {
96
0
                    if (op->swizzle.in[i] == op->swizzle.in[j])
97
0
                        return AVERROR(ENOTSUP);
98
0
                }
99
0
                p.index[i] = orig.index[op->swizzle.in[i]];
100
0
            }
101
0
            break;
102
0
        }
103
104
0
        case SWS_OP_CLEAR:
105
0
            for (int i = 0; i < 4; i++) {
106
0
                if (!SWS_COMP_TEST(op->clear.mask, i))
107
0
                    continue;
108
0
                if (op->clear.value[i].den != 1)
109
0
                    return AVERROR(ENOTSUP);
110
111
                /* Ensure all bytes to be cleared are the same, because we
112
                 * can't memset on multi-byte sequences */
113
0
                uint8_t val = op->clear.value[i].num & 0xFF;
114
0
                uint32_t ref = val;
115
0
                switch (ff_sws_pixel_type_size(op->type)) {
116
0
                case 2: ref *= 0x101; break;
117
0
                case 4: ref *= 0x1010101; break;
118
0
                }
119
0
                if (ref != op->clear.value[i].num)
120
0
                    return AVERROR(ENOTSUP);
121
0
                p.clear_value[i] = val;
122
0
                p.index[i] = -1;
123
0
            }
124
0
            break;
125
126
0
        case SWS_OP_WRITE:
127
0
            if (ff_sws_rw_op_planes(op) != op->rw.elems || op->rw.frac || op->rw.filter.op)
128
0
                return AVERROR(ENOTSUP);
129
0
            p.num_planes = op->rw.elems;
130
0
            break;
131
132
0
        default:
133
0
            return AVERROR(ENOTSUP);
134
0
        }
135
0
    }
136
137
0
    *out = (SwsCompiledOp) {
138
0
        .slice_align = 1,
139
0
        .block_size  = 1,
140
0
        .func = process,
141
0
        .priv = av_memdup(&p, sizeof(p)),
142
0
        .free = av_free,
143
0
    };
144
0
    return out->priv ? 0 : AVERROR(ENOMEM);
145
0
}
146
147
const SwsOpBackend backend_murder = {
148
    .name       = "memcpy",
149
    .flags      = SWS_BACKEND_MEMCPY,
150
    .compile    = compile,
151
    .hw_format  = AV_PIX_FMT_NONE,
152
};