Coverage Report

Created: 2026-06-17 15:31

next uncovered line (L), next uncovered region (R), next uncovered branch (B)
/tmp/bitcoin/src/univalue/lib/univalue_read.cpp
Line
Count
Source
1
// Copyright 2014 BitPay Inc.
2
// Distributed under the MIT software license, see the accompanying
3
// file COPYING or https://opensource.org/licenses/mit-license.php.
4
5
#include <univalue.h>
6
#include <univalue_utffilter.h>
7
8
#include <cstdint>
9
#include <cstring>
10
#include <string>
11
#include <string_view>
12
#include <vector>
13
14
/*
15
 * According to stackexchange, the original json test suite wanted
16
 * to limit depth to 22.  Widely-deployed PHP bails at depth 512,
17
 * so we will follow PHP's lead, which should be more than sufficient
18
 * (further stackexchange comments indicate depth > 32 rarely occurs).
19
 */
20
static constexpr size_t MAX_JSON_DEPTH = 512;
21
22
static bool json_isdigit(int ch)
23
7.41M
{
24
7.41M
    return ((ch >= '0') && (ch <= '9'));
25
7.41M
}
26
27
// convert hexadecimal string to unsigned integer
28
static const char *hatoui(const char *first, const char *last,
29
                          unsigned int& out)
30
592
{
31
592
    unsigned int result = 0;
32
2.96k
    for (; first != last; ++first)
33
2.36k
    {
34
2.36k
        int digit;
35
2.36k
        if (json_isdigit(*first))
36
1.17k
            digit = *first - '0';
37
38
1.19k
        else if (*first >= 'a' && *first <= 'f')
39
1.16k
            digit = *first - 'a' + 10;
40
41
22
        else if (*first >= 'A' && *first <= 'F')
42
22
            digit = *first - 'A' + 10;
43
44
0
        else
45
0
            break;
46
47
2.36k
        result = 16 * result + digit;
48
2.36k
    }
49
592
    out = result;
50
51
592
    return first;
52
592
}
53
54
enum jtokentype getJsonToken(std::string& tokenVal, unsigned int& consumed,
55
                            const char *raw, const char *end)
56
6.53M
{
57
6.53M
    tokenVal.clear();
58
6.53M
    consumed = 0;
59
60
6.53M
    const char *rawStart = raw;
61
62
8.40M
    while (raw < end && (json_isspace(*raw)))          // skip whitespace
63
1.86M
        raw++;
64
65
6.53M
    if (raw >= end)
66
181k
        return JTOK_NONE;
67
68
6.35M
    switch (*raw) {
69
70
318k
    case '{':
71
318k
        raw++;
72
318k
        consumed = (raw - rawStart);
73
318k
        return JTOK_OBJ_OPEN;
74
318k
    case '}':
75
318k
        raw++;
76
318k
        consumed = (raw - rawStart);
77
318k
        return JTOK_OBJ_CLOSE;
78
116k
    case '[':
79
116k
        raw++;
80
116k
        consumed = (raw - rawStart);
81
116k
        return JTOK_ARR_OPEN;
82
115k
    case ']':
83
115k
        raw++;
84
115k
        consumed = (raw - rawStart);
85
115k
        return JTOK_ARR_CLOSE;
86
87
928k
    case ':':
88
928k
        raw++;
89
928k
        consumed = (raw - rawStart);
90
928k
        return JTOK_COLON;
91
900k
    case ',':
92
900k
        raw++;
93
900k
        consumed = (raw - rawStart);
94
900k
        return JTOK_COMMA;
95
96
777
    case 'n':
97
16.9k
    case 't':
98
19.1k
    case 'f':
99
19.1k
        if (!strncmp(raw, "null", 4)) {
100
775
            raw += 4;
101
775
            consumed = (raw - rawStart);
102
775
            return JTOK_KW_NULL;
103
18.4k
        } else if (!strncmp(raw, "true", 4)) {
104
16.1k
            raw += 4;
105
16.1k
            consumed = (raw - rawStart);
106
16.1k
            return JTOK_KW_TRUE;
107
16.1k
        } else if (!strncmp(raw, "false", 5)) {
108
2.25k
            raw += 5;
109
2.25k
            consumed = (raw - rawStart);
110
2.25k
            return JTOK_KW_FALSE;
111
2.25k
        } else
112
7
            return JTOK_ERR;
113
114
40.5k
    case '-':
115
427k
    case '0':
116
957k
    case '1':
117
1.28M
    case '2':
118
1.47M
    case '3':
119
1.61M
    case '4':
120
1.67M
    case '5':
121
1.76M
    case '6':
122
1.82M
    case '7':
123
1.93M
    case '8':
124
2.00M
    case '9': {
125
        // part 1: int
126
2.00M
        std::string numStr;
127
128
2.00M
        const char *first = raw;
129
130
2.00M
        const char *firstDigit = first;
131
2.00M
        if (!json_isdigit(*firstDigit))
132
40.5k
            firstDigit++;
133
2.00M
        if ((*firstDigit == '0') && json_isdigit(firstDigit[1]))
134
1
            return JTOK_ERR;
135
136
2.00M
        numStr += *raw;                       // copy first char
137
2.00M
        raw++;
138
139
2.00M
        if ((*first == '-') && (raw < end) && (!json_isdigit(*raw)))
140
0
            return JTOK_ERR;
141
142
5.95M
        while (raw < end && json_isdigit(*raw)) {  // copy digits
143
3.95M
            numStr += *raw;
144
3.95M
            raw++;
145
3.95M
        }
146
147
        // part 2: frac
148
2.00M
        if (raw < end && *raw == '.') {
149
54.7k
            numStr += *raw;                   // copy .
150
54.7k
            raw++;
151
152
54.7k
            if (raw >= end || !json_isdigit(*raw))
153
0
                return JTOK_ERR;
154
586k
            while (raw < end && json_isdigit(*raw)) { // copy digits
155
532k
                numStr += *raw;
156
532k
                raw++;
157
532k
            }
158
54.7k
        }
159
160
        // part 3: exp
161
2.00M
        if (raw < end && (*raw == 'e' || *raw == 'E')) {
162
24.7k
            numStr += *raw;                   // copy E
163
24.7k
            raw++;
164
165
24.7k
            if (raw < end && (*raw == '-' || *raw == '+')) { // copy +/-
166
24.7k
                numStr += *raw;
167
24.7k
                raw++;
168
24.7k
            }
169
170
24.7k
            if (raw >= end || !json_isdigit(*raw))
171
3
                return JTOK_ERR;
172
74.2k
            while (raw < end && json_isdigit(*raw)) { // copy digits
173
49.5k
                numStr += *raw;
174
49.5k
                raw++;
175
49.5k
            }
176
24.7k
        }
177
178
2.00M
        tokenVal = numStr;
179
2.00M
        consumed = (raw - rawStart);
180
2.00M
        return JTOK_NUMBER;
181
2.00M
        }
182
183
1.63M
    case '"': {
184
1.63M
        raw++;                                // skip "
185
186
1.63M
        std::string valStr;
187
1.63M
        JSONUTF8StringFilter writer(valStr);
188
189
657M
        while (true) {
190
657M
            if (raw >= end || (unsigned char)*raw < 0x20)
191
4
                return JTOK_ERR;
192
193
657M
            else if (*raw == '\\') {
194
641
                raw++;                        // skip backslash
195
196
641
                if (raw >= end)
197
0
                    return JTOK_ERR;
198
199
641
                switch (*raw) {
200
16
                case '"':  writer.push_back('\"'); break;
201
4
                case '\\': writer.push_back('\\'); break;
202
2
                case '/':  writer.push_back('/'); break;
203
4
                case 'b':  writer.push_back('\b'); break;
204
4
                case 'f':  writer.push_back('\f'); break;
205
7
                case 'n':  writer.push_back('\n'); break;
206
4
                case 'r':  writer.push_back('\r'); break;
207
4
                case 't':  writer.push_back('\t'); break;
208
209
592
                case 'u': {
210
592
                    unsigned int codepoint;
211
592
                    if (raw + 1 + 4 >= end ||
212
592
                        hatoui(raw + 1, raw + 1 + 4, codepoint) !=
213
592
                               raw + 1 + 4)
214
0
                        return JTOK_ERR;
215
592
                    writer.push_back_u(codepoint);
216
592
                    raw += 4;
217
592
                    break;
218
592
                    }
219
4
                default:
220
4
                    return JTOK_ERR;
221
222
641
                }
223
224
637
                raw++;                        // skip esc'd char
225
637
            }
226
227
657M
            else if (*raw == '"') {
228
1.63M
                raw++;                        // skip "
229
1.63M
                break;                        // stop scanning
230
1.63M
            }
231
232
655M
            else {
233
655M
                writer.push_back(static_cast<unsigned char>(*raw));
234
655M
                raw++;
235
655M
            }
236
657M
        }
237
238
1.63M
        if (!writer.finalize())
239
4
            return JTOK_ERR;
240
1.63M
        tokenVal = valStr;
241
1.63M
        consumed = (raw - rawStart);
242
1.63M
        return JTOK_STRING;
243
1.63M
        }
244
245
26
    default:
246
26
        return JTOK_ERR;
247
6.35M
    }
248
6.35M
}
249
250
enum expect_bits : unsigned {
251
    EXP_OBJ_NAME = (1U << 0),
252
    EXP_COLON = (1U << 1),
253
    EXP_ARR_VALUE = (1U << 2),
254
    EXP_VALUE = (1U << 3),
255
    EXP_NOT_VALUE = (1U << 4),
256
};
257
258
23.2M
#define expect(bit) (expectMask & (EXP_##bit))
259
5.56M
#define setExpect(bit) (expectMask |= EXP_##bit)
260
5.75M
#define clearExpect(bit) (expectMask &= ~EXP_##bit)
261
262
bool UniValue::read(std::string_view str_in)
263
181k
{
264
181k
    clear();
265
266
181k
    uint32_t expectMask = 0;
267
181k
    std::vector<UniValue*> stack;
268
269
181k
    std::string tokenVal;
270
181k
    unsigned int consumed;
271
181k
    enum jtokentype tok = JTOK_NONE;
272
181k
    enum jtokentype last_tok = JTOK_NONE;
273
181k
    const char* raw{str_in.data()};
274
181k
    const char* end{raw + str_in.size()};
275
4.63M
    do {
276
4.63M
        last_tok = tok;
277
278
4.63M
        tok = getJsonToken(tokenVal, consumed, raw, end);
279
4.63M
        if (tok == JTOK_NONE || tok == JTOK_ERR)
280
47
            return false;
281
4.63M
        raw += consumed;
282
283
4.63M
        bool isValueOpen = jsonTokenIsValue(tok) ||
284
4.63M
            tok == JTOK_OBJ_OPEN || tok == JTOK_ARR_OPEN;
285
286
4.63M
        if (expect(VALUE)) {
287
928k
            if (!isValueOpen)
288
2
                return false;
289
928k
            clearExpect(VALUE);
290
291
3.71M
        } else if (expect(ARR_VALUE)) {
292
341k
            bool isArrValue = isValueOpen || (tok == JTOK_ARR_CLOSE);
293
341k
            if (!isArrValue)
294
2
                return false;
295
296
341k
            clearExpect(ARR_VALUE);
297
298
3.36M
        } else if (expect(OBJ_NAME)) {
299
994k
            bool isObjName = (tok == JTOK_OBJ_CLOSE || tok == JTOK_STRING);
300
994k
            if (!isObjName)
301
4
                return false;
302
303
2.37M
        } else if (expect(COLON)) {
304
928k
            if (tok != JTOK_COLON)
305
3
                return false;
306
928k
            clearExpect(COLON);
307
308
1.44M
        } else if (!expect(COLON) && (tok == JTOK_COLON)) {
309
1
            return false;
310
1
        }
311
312
4.63M
        if (expect(NOT_VALUE)) {
313
2.19M
            if (isValueOpen)
314
2
                return false;
315
2.19M
            clearExpect(NOT_VALUE);
316
2.19M
        }
317
318
4.63M
        switch (tok) {
319
320
318k
        case JTOK_OBJ_OPEN:
321
435k
        case JTOK_ARR_OPEN: {
322
435k
            VType utyp = (tok == JTOK_OBJ_OPEN ? VOBJ : VARR);
323
435k
            if (!stack.size()) {
324
181k
                if (utyp == VOBJ)
325
180k
                    setObject();
326
613
                else
327
613
                    setArray();
328
181k
                stack.push_back(this);
329
253k
            } else {
330
253k
                UniValue tmpVal(utyp);
331
253k
                UniValue *top = stack.back();
332
253k
                top->values.push_back(tmpVal);
333
334
253k
                UniValue *newTop = &(top->values.back());
335
253k
                stack.push_back(newTop);
336
253k
            }
337
338
435k
            if (stack.size() > MAX_JSON_DEPTH)
339
2
                return false;
340
341
435k
            if (utyp == VOBJ)
342
318k
                setExpect(OBJ_NAME);
343
116k
            else
344
116k
                setExpect(ARR_VALUE);
345
435k
            break;
346
435k
            }
347
348
318k
        case JTOK_OBJ_CLOSE:
349
434k
        case JTOK_ARR_CLOSE: {
350
434k
            if (!stack.size() || (last_tok == JTOK_COMMA))
351
2
                return false;
352
353
434k
            VType utyp = (tok == JTOK_OBJ_CLOSE ? VOBJ : VARR);
354
434k
            UniValue *top = stack.back();
355
434k
            if (utyp != top->getType())
356
1
                return false;
357
358
434k
            stack.pop_back();
359
434k
            clearExpect(OBJ_NAME);
360
434k
            setExpect(NOT_VALUE);
361
434k
            break;
362
434k
            }
363
364
928k
        case JTOK_COLON: {
365
928k
            if (!stack.size())
366
0
                return false;
367
368
928k
            UniValue *top = stack.back();
369
928k
            if (top->getType() != VOBJ)
370
0
                return false;
371
372
928k
            setExpect(VALUE);
373
928k
            break;
374
928k
            }
375
376
900k
        case JTOK_COMMA: {
377
900k
            if (!stack.size() ||
378
900k
                (last_tok == JTOK_COMMA) || (last_tok == JTOK_ARR_OPEN))
379
0
                return false;
380
381
900k
            UniValue *top = stack.back();
382
900k
            if (top->getType() == VOBJ)
383
675k
                setExpect(OBJ_NAME);
384
224k
            else
385
224k
                setExpect(ARR_VALUE);
386
900k
            break;
387
900k
            }
388
389
774
        case JTOK_KW_NULL:
390
16.9k
        case JTOK_KW_TRUE:
391
19.1k
        case JTOK_KW_FALSE: {
392
19.1k
            UniValue tmpVal;
393
19.1k
            switch (tok) {
394
774
            case JTOK_KW_NULL:
395
                // do nothing more
396
774
                break;
397
16.1k
            case JTOK_KW_TRUE:
398
16.1k
                tmpVal.setBool(true);
399
16.1k
                break;
400
2.25k
            case JTOK_KW_FALSE:
401
2.25k
                tmpVal.setBool(false);
402
2.25k
                break;
403
0
            default: /* impossible */ break;
404
19.1k
            }
405
406
19.1k
            if (!stack.size()) {
407
39
                *this = tmpVal;
408
39
                break;
409
39
            }
410
411
19.1k
            UniValue *top = stack.back();
412
19.1k
            top->values.push_back(tmpVal);
413
414
19.1k
            setExpect(NOT_VALUE);
415
19.1k
            break;
416
19.1k
            }
417
418
288k
        case JTOK_NUMBER: {
419
288k
            UniValue tmpVal(VNUM, tokenVal);
420
288k
            if (!stack.size()) {
421
69
                *this = tmpVal;
422
69
                break;
423
69
            }
424
425
288k
            UniValue *top = stack.back();
426
288k
            top->values.push_back(tmpVal);
427
428
288k
            setExpect(NOT_VALUE);
429
288k
            break;
430
288k
            }
431
432
1.63M
        case JTOK_STRING: {
433
1.63M
            if (expect(OBJ_NAME)) {
434
928k
                UniValue *top = stack.back();
435
928k
                top->keys.push_back(tokenVal);
436
928k
                clearExpect(OBJ_NAME);
437
928k
                setExpect(COLON);
438
928k
            } else {
439
705k
                UniValue tmpVal(VSTR, tokenVal);
440
705k
                if (!stack.size()) {
441
3
                    *this = tmpVal;
442
3
                    break;
443
3
                }
444
705k
                UniValue *top = stack.back();
445
705k
                top->values.push_back(tmpVal);
446
705k
            }
447
448
1.63M
            setExpect(NOT_VALUE);
449
1.63M
            break;
450
1.63M
            }
451
452
0
        default:
453
0
            return false;
454
4.63M
        }
455
4.63M
    } while (!stack.empty ());
456
457
    /* Check that nothing follows the initial construct (parsed above).  */
458
181k
    tok = getJsonToken(tokenVal, consumed, raw, end);
459
181k
    if (tok != JTOK_NONE)
460
13
        return false;
461
462
181k
    return true;
463
181k
}
464