X-Git-Url: https://pintos-os.org/cgi-bin/gitweb.cgi?a=blobdiff_plain;f=dump.c;h=cba4509cd41ebdecdff224aba3c46e5a7ba374b3;hb=92a04dcbe26bed7d6c044f06414bc54e0dbb1ad1;hp=f2eaa5aec1082a94d904657ff8ebfa1b984cb46d;hpb=e2f2f98013f9dc4867e55fe4b8b551a595726886;p=pspp diff --git a/dump.c b/dump.c index f2eaa5aec1..cba4509cd4 100644 --- a/dump.c +++ b/dump.c @@ -1,3 +1,4 @@ +#include #include #include #include @@ -9,6 +10,8 @@ static uint8_t *data; static size_t n; +int version; + static bool all_ascii(const uint8_t *p, size_t n) { @@ -25,13 +28,6 @@ try_find(const char *target, size_t target_len) return pos ? pos - data : 0; } -static size_t -try_find_tail(const char *target, size_t target_len) -{ - size_t pos = try_find(target, target_len); - return pos ? pos + target_len : 0; -} - static size_t find(const char *target, size_t target_len) { @@ -44,18 +40,6 @@ find(const char *target, size_t target_len) return pos; } -static size_t -find_tail(const char *target, size_t target_len) -{ - size_t pos = try_find_tail(target, target_len); - if (!pos) - { - fprintf (stderr, "not found\n"); - exit(1); - } - return pos; -} - size_t pos; #define XSTR(x) #x @@ -124,870 +108,340 @@ match_byte_assert(uint8_t b, const char *where) } #define match_byte_assert(b) match_byte_assert(b, WHERE) -static char * -get_string(const char *where) +static void +newline(FILE *stream, int pos) { - if (1 - /*data[pos + 1] == 0 && data[pos + 2] == 0 && data[pos + 3] == 0*/ - /*&& all_ascii(&data[pos + 4], data[pos])*/) - { - int len = data[pos] + data[pos + 1] * 256; - char *s = malloc(len + 1); - - memcpy(s, &data[pos + 4], len); - s[len] = 0; - pos += 4 + len; - return s; - } - else - { - fprintf(stderr, "%s: 0x%x: expected string\n", where, pos); - exit(1); - } + fprintf(stream, "\n%08x: ", pos); } -#define get_string() get_string(WHERE) static void -dump_value_31(void) +dump_raw(FILE *stream, int start, int end) { - if (match_byte (0x31)) + for (size_t i = start; i < end; ) { - if (match_u32 (0)) + if (i + 5 <= n + && data[i] + //&& !data[i + 1] + && !data[i + 2] + && !data[i + 3] + && i + 4 + data[i] + data[i + 1] * 256 <= end + && all_ascii(&data[i + 4], data[i] + data[i + 1] * 256)) { - match_u32_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; + newline(stream, i); + fprintf(stream, "\""); + fwrite(&data[i + 4], 1, data[i] + data[i + 1] * 256, stream); + fputs("\" ", stream); + + i += 4 + data[i] + data[i + 1] * 256; } - else if (match_u32 (1)) + else if (i + 12 <= end + && data[i + 1] == 40 + && data[i + 2] == 5 + && data[i + 3] == 0) { - printf("(footnote %d) ", get_u32()); - match_byte_assert (0); - match_byte_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; + double d; + + memcpy (&d, &data[i + 4], 8); + fprintf (stream, "F40.%d(%.*f)", data[i], data[i], d); + i += 12; + newline (stream, i); } - else if (match_u32 (2)) + else if (i + 12 <= end + && data[i + 1] == 40 + && data[i + 2] == 31 + && data[i + 3] == 0) { - printf("(special 2)"); - match_byte_assert(0); - match_byte_assert(0); - match_u32_assert(1); - match_byte_assert(0); - match_byte_assert(0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; + double d; + + memcpy (&d, &data[i + 4], 8); + fprintf (stream, "PCT40.%d(%.*f)", data[i], data[i], d); + i += 12; + newline(stream, i); + } + else if (i + 4 <= end + && (data[i] && data[i] != 88 && data[i] != 0x41) + && !data[i + 1] + && !data[i + 2] + && !data[i + 3]) + { + fprintf (stream, "i%d ", data[i]); + i += 4; } else { - match_u32_assert(3); - printf("(special 3)"); - match_byte_assert(0); - match_byte_assert(0); - match_byte_assert(1); - match_byte_assert(0); - int subn = get_u32 (); - printf ("nested %d bytes, ", subn); - pos += subn; - subn = get_u32 (); - printf ("nested %d bytes, ", subn); - pos += subn; + fprintf(stream, "%02x ", data[i]); + i++; } } - else - match_byte_assert (0x58); -} -static void -dump_value(int level) -{ - for (int i = 0; i <= level; i++) - printf (" "); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - if (match_byte (3)) - { - char *s1 = get_string(); - dump_value_31(); - char *s2 = get_string(); - char *s3 = get_string(); - if (strcmp(s1, s3)) - printf("strings \"%s\", \"%s\" and \"%s\"", s1, s2, s3); - else - printf("string \"%s\" and \"%s\"", s1, s2); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (1); - match_byte (1); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (1); - } - else if (match_byte (5)) - { - match_byte_assert (0x58); - printf ("variable \"%s\"", get_string()); - get_string(); - if (!match_byte(1) && !match_byte(2)) - match_byte_assert(3); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - } - else if (match_byte (2)) - { - unsigned int format; - char *var, *vallab; - double value; - - match_byte_assert (0x58); - format = get_u32 (); - value = get_double (); - var = get_string (); - vallab = get_string (); - printf ("value %g format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); - if (!match_byte (1) && !match_byte(2)) - match_byte_assert (3); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - } - else if (match_byte (4)) - { - unsigned int format; - char *var, *vallab, *value; +} - match_byte_assert (0x58); - format = get_u32 (); - vallab = get_string (); - var = get_string (); - if (!match_byte(1) && !match_byte(2)) - match_byte_assert (3); - value = get_string (); - printf ("value \"%s\" format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - } - else if (match_byte (1)) +static char * +get_string(const char *where) +{ + if (1 + /*data[pos + 1] == 0 && data[pos + 2] == 0 && data[pos + 3] == 0*/ + /*&& all_ascii(&data[pos + 4], data[pos])*/) { - unsigned int format; - double value; + int len = data[pos] + data[pos + 1] * 256; + char *s = malloc(len + 1); - dump_value_31(); - format = get_u32 (); - value = get_double (); - printf ("value %g format %d(%d.%d)", value, format >> 16, (format >> 8) & 0xff, format & 0xff); - match_byte (1); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (1); + memcpy(s, &data[pos + 4], len); + s[len] = 0; + pos += 4 + len; + return s; } else { - dump_value_31(); - char *base = get_string(); - int x = get_u32(); - printf ("\"%s\" with %d variables:\n", base, x); - if (match_u32(0)) - { - for (int i = 0; i < x; i++) - { - dump_value (level+1); - putchar('\n'); - } - } - else - { - for (int i = 0; i < x; i++) - { - int y = get_u32(); - match_u32_assert(0); - for (int j = 0; j <= level; j++) - printf (" "); - printf("variable %d has %d values:\n", i, y); - for (int j = 0; j < y; j++) - { - if (match_byte(3)) - { - char *a = get_string(); - match_byte_assert(0x58); - char *b = get_string(); - char *c = get_string(); - for (int k = 0; k <= level + 1; k++) - printf (" "); - printf ("\"%s\", \"%s\", \"%s\"", a, b, c); - match_byte(0); - match_byte(0); - match_byte(0); - match_byte(0); - match_byte(0); - } - else - dump_value (level+1); - putchar('\n'); - } - } - } + fprintf(stderr, "%s: 0x%x: expected string\n", where, pos); + exit(1); } } +#define get_string() get_string(WHERE) -static void -dump_dim_value(int level) +static char * +dump_nested_string(void) { - for (int i = 0; i <= level; i++) - printf (" "); + char *s = NULL; - if (match_byte (3)) - { - get_string(); - if (match_byte (0x31)) - { - match_u32 (1); - printf("(footnote %d) ", get_u32()); - match_byte_assert (0); - match_byte_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - else - match_byte_assert (0x58); - get_string(); - printf("string \"%s\"", get_string()); - match_byte (0); - match_byte (1); - } - else if (match_byte (5)) + match_byte_assert (0); + match_byte_assert (0); + int outer_end = pos + get_u32(); + int inner_end = pos + get_u32(); + if (pos != inner_end) { - match_byte_assert (0x58); - printf ("variable \"%s\"", get_string()); - get_string(); - if (!match_byte (2)) - match_byte_assert(3); - } - else if (match_byte (2)) - { - unsigned int format; - char *var, *vallab; - double value; - - match_byte_assert (0x58); - format = get_u32 (); - value = get_double (); - var = get_string (); - vallab = get_string (); - printf ("value %g format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); - if (!match_u32 (3)) - match_u32_assert (2); - match_byte (0); - } - else if (match_byte (1)) - { - unsigned int format; - double value; - - match_byte_assert (0x58); - format = get_u32 (); - value = get_double (); - printf ("value %g format %d(%d.%d)", value, format >> 16, (format >> 8) & 0xff, format & 0xff); - match_byte (1); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (1); - } - else - { - int subn; - - match_byte (0); - if (match_byte (0x31)) - { - match_u32_assert (0); - match_u32_assert (0); - subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } + match_u32_assert(0); + if (match_byte(0x31)) + s = get_string(); else match_byte_assert(0x58); - printf ("; \"%s\", substitutions:", get_string()); - int total_subs = get_u32(); - int x = get_u32(); - if (x) + if (pos != inner_end) { - total_subs = (total_subs - 1) + x; - match_u32_assert (0); - } - printf (" (total %d)", total_subs); - - for (int i = 0; i < total_subs; i++) - { - putc ('\n', stdout); - dump_value (level + 1); + fprintf(stderr, "inner end discrepancy\n"); + exit(1); } } -} - -static void -dump_category(int level) -{ - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - dump_value (level); - - if (match_u32 (2)) - get_u32 (); - else if (match_u32 (1)) - { - match_byte (0); - match_byte (0); - match_byte (0); - get_u32 (); - } - else if (match_byte (1)) - { - match_byte (0); - if (!match_u32 (2)) - match_u32_assert (1); - match_byte (0); - get_u32(); - } - else + match_byte_assert(0x58); + match_byte_assert(0x58); + if (pos != outer_end) { - match_u32_assert (0); - get_u32 (); + fprintf(stderr, "outer end discrepancy\n"); + exit(1); } - int n_categories = get_u32(); - if (n_categories > 0) - printf (", %d subcategories:", n_categories); - printf("\n"); - for (int i = 0; i < n_categories; i++) - dump_category (level + 1); -} - -static void -dump_dim(void) -{ - int n_categories; - printf("next dim\n"); - match_byte(0); - dump_dim_value(0); - - /* This byte is usually 0x02 but 0x00 and 0x75 (!) have also been spotted. */ - pos++; - - if (!match_byte(0) && !match_byte(1)) - match_byte_assert(2); - if (!match_u32(0)) - match_u32_assert(2); - if (!match_byte(0)) - match_byte_assert(1); - match_byte(0); - match_byte(0); - match_byte(0); - match_byte(0); - get_u32(); - match_byte(0); - match_byte(0); - match_byte(0); - match_byte(0); - n_categories = get_u32(); - printf("%d nested categories\n", n_categories); - for (int i = 0; i < n_categories; i++) - dump_category (0); -} - -int n_dims; -static void -dump_dims(void) -{ - n_dims = get_u32(); - printf ("%u dimensions\n", n_dims); - for (int i = 0; i < n_dims; i++) - { - printf("\n"); - dump_dim (); - } + return s; } static void -dump_data_value_31(void) +dump_value_31(FILE *stream) { if (match_byte (0x31)) { if (match_u32 (0)) { if (match_u32 (1)) - get_string(); + { + /* Only "a" observed as a sample value (although it appears 44 times in the corpus). */ + get_string(); + } else match_u32_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - else if (match_u32 (1)) - { - printf("(footnote %d) ", get_u32()); - match_byte_assert (0); - match_byte_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - else if (match_u32 (2)) - { - printf("(special 2)"); - match_byte_assert(0); - match_byte_assert(0); - match_u32_assert(1); - match_byte_assert(0); - match_byte_assert(0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - else - { - match_u32_assert(3); - printf("(special 3)"); - match_byte_assert(0); - match_byte_assert(0); - match_byte_assert(1); - match_byte_assert(0); - int subn = get_u32 (); - printf ("nested %d bytes, ", subn); - pos += subn; - subn = get_u32 (); - printf ("nested %d bytes, ", subn); - pos += subn; - } - } - else - match_byte_assert (0x58); -} -static void -dump_data(void) -{ -#if 1 - int a[16]; - for (int i = 0; i < 3 + n_dims; i++) - a[i] = get_u32(); - printf ("data intro:"); - for (int i = 0; i < 3 + n_dims; i++) - printf(" %d", a[i]); - printf("\n"); -#else - fprintf (stderr,"data intro (%d dims):", n_dims); - for (int i = 0; i < 3+n_dims; i++) - fprintf (stderr," %d", get_u32()); - fprintf(stderr,"\n"); -#endif - int x = get_u32(); - printf ("%d data values, starting at %08x\n", x, pos); - for (int i = 0; i < x; i++) - { - printf("%08x, index %d:\n", pos, get_u32()); - match_u32_assert(0); - match_byte(0); - match_byte(0); - match_byte(0); - match_byte(0); - if (match_byte (1)) - { - unsigned int format; - double value; + if (version == 1) + { + /* We only have one SPV file for this version (with many + tables). */ + match_byte(0); + if (!match_u32(1)) + match_u32_assert(2); + match_byte(0); + match_byte(0); + if (!match_u32(0) && !match_u32(1) && !match_u32(2) && !match_u32(3) && !match_u32(4) && !match_u32(5) && !match_u32(6) && !match_u32(7) && !match_u32(8) && !match_u32(9)) + match_u32_assert(10); + match_byte(0); + match_byte(0); + return; + } - dump_data_value_31(); - format = get_u32 (); - value = get_double (); - printf (" value %g format %d(%d.%d)", value, format >> 16, (format >> 8) & 0xff, format & 0xff); - } - else if (match_byte (3)) - { - get_string(); - dump_data_value_31(); - get_string(); - printf("string \"%s\"", get_string()); - match_byte (0); - } - else if (match_byte (2)) - { - unsigned int format; - char *var, *vallab; - double value; - - match_byte_assert (0x58); - format = get_u32 (); - value = get_double (); - var = get_string (); - vallab = get_string (); - printf ("value %g format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); - if (!match_byte (1) && !match_byte(2)) - match_byte_assert (3); - } - else if (match_byte (4)) - { - unsigned int format; - char *var, *vallab, *value; - - match_byte_assert (0x58); - format = get_u32 (); - vallab = get_string (); - var = get_string (); - if (!match_byte(1) && !match_byte(2)) - match_byte_assert (3); - value = get_string (); - printf ("value \"%s\" format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); - } - else if (match_byte (5)) - { - match_byte_assert (0x58); - printf ("variable \"%s\"", get_string()); - get_string(); - if (!match_byte(1) && !match_byte(2)) - match_byte_assert(3); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - } - else - { - dump_data_value_31(); - char *base = get_string(); - int x = get_u32(); - printf ("\"%s\"; %d variables:\n", base, x); - for (int i = 0; i < x; i++) + int outer_end = pos + get_u32(); + int inner_end = pos + get_u32(); + if (pos != inner_end) { - int y = get_u32(); - if (!y) - y = 1; + match_u32_assert(0); + if (match_byte(0x31)) + { + /* Appears to be a template string, e.g. '^1 cells (^2) expf < 5. Min exp = ^3...'. + Probably doesn't actually appear in output because many examples look unpolished, + e.g. 'partial list cases value ^1 shown upper...' */ + get_string(); + } else - match_u32_assert(0); - for (int j = 0; j <= 0; j++) - printf (" "); - printf("variable %d has %d values:\n", i, y); - for (int j = 0; j < y; j++) + match_byte_assert(0x58); + if (pos != inner_end) { - if (match_byte (1)) - { - unsigned int format; - double value; - - dump_data_value_31(); - format = get_u32 (); - value = get_double (); - printf (" value %g format %d(%d.%d)", value, format >> 16, (format >> 8) & 0xff, format & 0xff); - } - else if (match_byte(3)) - { - char *a = get_string(); - match_byte_assert(0x58); - char *b = get_string(); - char *c = get_string(); - for (int k = 0; k <= 1; k++) - printf (" "); - printf ("\"%s\", \"%s\", \"%s\"", a, b, c); - match_byte(0); - match_byte(0); - match_byte(0); - match_byte(0); - } - else if (match_byte(5)) - { - match_byte_assert (0x58); - printf ("variable \"%s\"", get_string()); - get_string(); - if (!match_byte(1) && !match_byte(2)) - match_byte_assert(3); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - } - else - dump_value (0); - putchar('\n'); + fprintf(stderr, "inner end discrepancy\n"); + exit(1); } } - } - putchar('\n'); - } -} -static void -dump_title_value_31(int level) -{ - if (match_byte (0x31)) - { - if (match_u32 (0)) - { - match_u32_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; + if (match_byte(0x31)) + { + /* Only one example in the corpus. */ + match_byte(1); + match_byte(0); + match_byte(0); + match_byte(0); + match_byte_assert(1); + get_string(); /* foreground */ + get_string(); /* background */ + get_string(); /* font */ + if (!match_byte(14)) + match_byte_assert(12); /* size? */ + } + else + match_byte_assert(0x58); + if (match_byte(0x31)) + { + /* Only two SPV files have anything like this, so it's hard to + generalize. */ + match_u32_assert(0); + match_u32_assert(0); + match_u32_assert(0); + match_u32_assert(0); + match_byte_assert(1); + match_byte_assert(0); + if (!match_byte(8) && !match_byte(1)) + match_byte_assert(2); + match_byte_assert(0); + match_byte_assert(8); + match_byte_assert(0); + match_byte_assert(10); + match_byte_assert(0); + } + else + match_byte_assert(0x58); + if (pos != outer_end) + { + fprintf(stderr, "outer end discrepancy\n"); + exit(1); + } } else if (match_u32 (1)) { - printf("(footnote %d) ", get_u32()); - match_byte_assert (0); - match_byte_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; + fprintf(stream, "(footnote %d) ", get_u32()); + dump_nested_string(); } else if (match_u32 (2)) { - printf("(special 2)"); - match_byte_assert(0); - match_byte_assert(0); - if (!match_u32(2)) - match_u32_assert(1); - match_byte_assert(0); + fprintf(stream, "(special 2)"); + if (!match_byte(0)) + match_byte_assert(2); match_byte_assert(0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; + if (!match_u32 (2) && !match_u32(1)) + match_u32_assert(3); + dump_nested_string(); /* Our corpus doesn't contain any examples with strings though. */ } else { match_u32_assert(3); - printf("(special 3)"); + fprintf(stream, "(special 3)"); match_byte_assert(0); match_byte_assert(0); match_byte_assert(1); match_byte_assert(0); - int subn = get_u32 (); - printf ("nested %d bytes, ", subn); - pos += subn; - subn = get_u32 (); - printf ("nested %d bytes, ", subn); - pos += subn; + match_u32_assert(2); + dump_nested_string(); /* Our corpus doesn't contain any examples with strings though. */ } } else match_byte_assert (0x58); -} - -static void -dump_title_value(int level) -{ - for (int i = 0; i <= level; i++) - printf (" "); - - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - if (match_byte (3)) - { - get_string(); - dump_title_value_31(level); - get_string(); - printf("string \"%s\"", get_string()); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (1); - match_byte (1); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (1); - } - else if (match_byte (5)) - { - dump_title_value_31(level); - printf ("variable \"%s\"", get_string()); - get_string(); - if (!match_byte(1) && !match_byte(2)) - match_byte_assert(3); - } - else if (match_byte (2)) - { - unsigned int format; - char *var, *vallab; - double value; - - match_byte_assert (0x58); - format = get_u32 (); - value = get_double (); - var = get_string (); - vallab = get_string (); - printf ("value %g format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); - if (!match_byte (1) && !match_byte(2)) - match_byte_assert (3); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - } - else if (match_byte (4)) - { - unsigned int format; - char *var, *vallab, *value; - - match_byte_assert (0x58); - format = get_u32 (); - vallab = get_string (); - var = get_string (); - if (!match_byte(1) && !match_byte(2)) - match_byte_assert (3); - value = get_string (); - printf ("value \"%s\" format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - } - else if (match_byte (1)) - { - unsigned int format; - double value; - - dump_title_value_31(level); - format = get_u32 (); - value = get_double (); - printf ("value %g format %d(%d.%d)", value, format >> 16, (format >> 8) & 0xff, format & 0xff); - match_byte (1); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (1); - } - else - { - dump_title_value_31(level); - - char *base = get_string(); - int x = get_u32(); - printf ("\"%s\" with %d variables:\n", base, x); - for (int i = 0; i < x; i++) - { - int y = get_u32(); - if (!y) - y = 1; - else - match_u32_assert(0); - for (int j = 0; j <= level; j++) - printf (" "); - printf("variable %d has %d values:\n", i, y); - for (int j = 0; j < y; j++) - { - match_byte(0); - if (match_byte(3)) - { - char *a = get_string(); - match_byte_assert(0x58); - char *b = get_string(); - char *c = get_string(); - for (int k = 0; k <= level + 1; k++) - printf (" "); - printf ("\"%s\", \"%s\", \"%s\"", a, b, c); - } - else - dump_title_value (level+1); - putchar('\n'); - } - } +} + +static const char * +format_to_string (int type) +{ + static char tmp[16]; + switch (type) + { + case 1: return "A"; + case 2: return "AHEX"; + case 3: return "COMMA"; + case 4: return "DOLLAR"; + case 5: case 40: return "F"; + case 6: return "IB"; + case 7: return "PIBHEX"; + case 8: return "P"; + case 9: return "PIB"; + case 10: return "PK"; + case 11: return "RB"; + case 12: return "RBHEX"; + case 15: return "Z"; + case 16: return "N"; + case 17: return "E"; + case 20: return "DATE"; + case 21: return "TIME"; + case 22: return "DATETIME"; + case 23: return "ADATE"; + case 24: return "JDATE"; + case 25: return "DTIME"; + case 26: return "WKDAY"; + case 27: return "MONTH"; + case 28: return "MOYR"; + case 29: return "QYR"; + case 30: return "WKYR"; + case 31: return "PCT"; + case 32: return "DOT"; + case 33: return "CCA"; + case 34: return "CCB"; + case 35: return "CCC"; + case 36: return "CCD"; + case 37: return "CCE"; + case 38: return "EDATE"; + case 39: return "SDATE"; + default: + abort(); + sprintf(tmp, "<%d>", type); + return tmp; } } static void -dump_footnote_value(int level) +dump_value(FILE *stream, int level) { + match_byte(0); + match_byte(0); + match_byte(0); + match_byte(0); + for (int i = 0; i <= level; i++) - printf (" "); + fprintf (stream, " "); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); if (match_byte (3)) { - get_string(); - if (match_byte (0x31)) - { - if (match_u32 (1)) - { - printf("(footnote %d) ", get_u32()); - match_byte_assert (0); - match_byte_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - else if (match_u32 (2)) - { - printf("(special 2)"); - match_byte_assert(0); - match_byte_assert(0); - match_u32_assert(1); - match_byte_assert(0); - match_byte_assert(0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - else - { - match_u32_assert(3); - printf("(special 3)"); - match_byte_assert(0); - match_byte_assert(0); - match_byte_assert(1); - match_byte_assert(0); - int subn = get_u32 (); - printf ("nested %d bytes, ", subn); - pos += subn; - subn = get_u32 (); - printf ("nested %d bytes, ", subn); - pos += subn; - } - } - else - match_byte_assert (0x58); - get_string(); - printf("string \"%s\"", get_string()); + char *text = get_string(); + dump_value_31(stream); + char *identifier = get_string(); + char *text_eng = get_string(); + fprintf (stream, "\n"); if (!match_byte (0)) - match_byte_assert (1); + match_byte_assert(1); + } else if (match_byte (5)) { - match_byte_assert (0x58); - printf ("variable \"%s\"", get_string()); - get_string(); + dump_value_31(stream); + char *name = get_string (); + char *label = get_string (); + fprintf (stream, "\n"); if (!match_byte(1) && !match_byte(2)) match_byte_assert(3); } @@ -997,22 +451,20 @@ dump_footnote_value(int level) char *var, *vallab; double value; - match_byte_assert (0x58); + dump_value_31 (stream); format = get_u32 (); value = get_double (); var = get_string (); vallab = get_string (); - printf ("value %g format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); + fprintf (stream, "> 16), (format >> 8) & 0xff, format & 0xff); + if (var[0]) + fprintf (stream, " variable=\"%s\"", var); + if (vallab[0]) + fprintf (stream, " label=\"%s\"/>\n", vallab); + fprintf (stream, "/>\n"); if (!match_byte (1) && !match_byte(2)) match_byte_assert (3); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); } else if (match_byte (4)) { @@ -1026,58 +478,32 @@ dump_footnote_value(int level) if (!match_byte(1) && !match_byte(2)) match_byte_assert (3); value = get_string (); - printf ("value \"%s\" format %d(%d.%d) var \"%s\" vallab \"%s\"", - value, format >> 16, (format >> 8) & 0xff, format & 0xff, var, vallab); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (0); + fprintf (stream, "> 16), (format >> 8) & 0xff, format & 0xff); + if (var[0]) + fprintf (stream, " variable=\"%s\"", var); + if (vallab[0]) + fprintf (stream, " label=\"%s\"/>\n", vallab); + fprintf (stream, "/>\n"); } else if (match_byte (1)) { unsigned int format; double value; - if (match_byte (0x31)) - { - if (match_u32 (1)) - { - printf("(footnote %d) ", get_u32()); - match_byte_assert (0); - match_byte_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - } - else - match_byte_assert (0x58); + dump_value_31(stream); format = get_u32 (); value = get_double (); - printf ("value %g format %d(%d.%d)", value, format >> 16, (format >> 8) & 0xff, format & 0xff); + fprintf (stream, "\n", + DBL_DIG, value, format_to_string(format >> 16), (format >> 8) & 0xff, format & 0xff); } - else if (match_byte (0x31)) + else { - if (match_u32 (1)) - { - printf("(footnote %d) ", get_u32()); - match_byte_assert (0); - match_byte_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - else - { - match_u32_assert (0); - match_u32_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } + dump_value_31(stream); + char *base = get_string(); int x = get_u32(); - printf ("\"%s\"; %d variables:\n", base, x); + fprintf (stream, "\n"); } - else +} + +static int +compare_int(const void *a_, const void *b_) +{ + const int *a = a_; + const int *b = b_; + return *a < *b ? -1 : *a > *b; +} + +static void +check_permutation(int *a, int n, const char *name) +{ + int b[n]; + memcpy(b, a, n * sizeof *a); + qsort(b, n, sizeof *b, compare_int); + for (int i = 0; i < n; i++) + if (b[i] != i) + { + fprintf(stderr, "bad %s permutation:", name); + for (int i = 0; i < n; i++) + fprintf(stderr, " %d", a[i]); + putc('\n', stderr); + exit(1); + } +} + +static void +dump_category(int level, int *indexes, int *n_indexes, int max_indexes) +{ + for (int i = 0; i <= level; i++) + fprintf (stdout, " "); + printf ("\n"); + dump_value (stdout, level + 1); + match_u32(1); + + if (match_u32 (1)) + match_byte (0); + else if (match_byte (1)) { + match_byte (0); + if (!match_u32 (2)) + match_u32_assert (1); + match_byte (0); + } + else if (!match_u32(2)) + match_u32_assert (0); - match_byte_assert (0x58); - char *base = get_string(); - int x = get_u32(); - printf ("\"%s\" with %d variables:\n", base, x); - for (int i = 0; i < x; i++) + int indx = get_u32(); + int n_categories = get_u32(); + if (indx != -1) + { + if (n_categories != 0) { - int y = get_u32(); - if (!y) - y = 1; - else - match_u32_assert(0); - for (int j = 0; j <= level; j++) - printf (" "); - printf("variable %d has %d values:\n", i, y); - for (int j = 0; j < y; j++) - { - if (match_byte(3)) - { - char *a = get_string(); - match_byte_assert(0x58); - char *b = get_string(); - char *c = get_string(); - for (int k = 0; k <= level + 1; k++) - printf (" "); - printf ("\"%s\", \"%s\", \"%s\"", a, b, c); - match_byte_assert(0); - } - else - dump_footnote_value (level+1); - putchar('\n'); - } + fprintf(stderr, "index not -1 but subcategories\n"); + exit(1); + } + if (*n_indexes >= max_indexes) + { + fprintf(stderr, "too many categories (increase max_indexes)\n"); + exit(1); } + indexes[(*n_indexes)++] = indx; + } + if (n_categories == 0) + { + for (int i = 0; i <= level + 1; i++) + fprintf (stdout, " "); + fprintf (stdout, "%d\n", indx); + } + for (int i = 0; i < n_categories; i++) + dump_category (level + 1, indexes, n_indexes, max_indexes); + for (int i = 0; i <= level; i++) + fprintf (stdout, " "); + printf ("\n"); +} + +static void +dump_dim(int indx) +{ + int n_categories; + + printf ("\n", indx); + dump_value (stdout, 0); + + /* This byte is usually 0x02 but many other values have been spotted. */ + pos++; + + if (!match_byte(0) && !match_byte(1)) + match_byte_assert(2); + if (!match_u32(0)) + match_u32_assert(2); + if (!match_byte(0)) + match_byte_assert(1); + if (!match_byte(0)) + match_byte_assert(1); + match_byte_assert(1); + if (!match_u32(UINT32_MAX)) + match_u32_assert(indx); + n_categories = get_u32(); + + int indexes[2048]; + int n_indexes = 0; + for (int i = 0; i < n_categories; i++) + dump_category (0, indexes, &n_indexes, sizeof indexes / sizeof *indexes); + check_permutation(indexes, n_indexes, "categories"); + + fprintf (stdout, "\n"); +} + +int n_dims; +static void +dump_dims(void) +{ + n_dims = get_u32(); + for (int i = 0; i < n_dims; i++) + dump_dim (i); +} + +static void +dump_data(void) +{ + /* The first three numbers add to the number of dimensions. */ + int t = get_u32(); + t += get_u32(); + match_u32_assert(n_dims - t); + + /* The next n_dims numbers are a permutation of the dimension numbers. */ + int a[n_dims]; + for (int i = 0; i < n_dims; i++) + a[i] = get_u32(); + check_permutation(a, n_dims, "dimensions"); + + int x = get_u32(); + printf ("\n"); + for (int i = 0; i < x; i++) + { + printf (" \n", get_u32()); + match_u32_assert(0); + if (version == 1) + match_byte(0); + dump_value(stdout, 1); + fprintf (stdout, " \n"); } + printf ("\n"); } static void dump_title(void) { pos = 0x27; - dump_title_value(0); putchar('\n'); - dump_title_value(0); putchar('\n'); + printf ("\n"); + dump_value(stdout, 0); + match_byte(1); + printf ("\n"); + + printf ("\n"); + dump_value(stdout, 0); + match_byte(1); + printf ("\n"); + match_byte_assert(0x31); - dump_title_value(0); putchar('\n'); + + printf ("\n"); + dump_value(stdout, 0); + match_byte(1); + printf ("\n"); + match_byte(0); match_byte_assert(0x58); if (match_byte(0x31)) { - dump_footnote_value(0); putchar('\n'); + printf ("\n"); + dump_value(stdout, 0); + printf ("\n"); } else match_byte_assert(0x58); int n_footnotes = get_u32(); - if (n_footnotes >= 20) - { - fprintf(stderr, "%08x: %d footnotes\n", pos - 4, n_footnotes); - exit(1); - } - - printf("------\n%d footnotes\n", n_footnotes); - if (n_footnotes < 20) + for (int i = 0; i < n_footnotes; i++) { - for (int i = 0; i < n_footnotes; i++) + printf ("\n", i); + dump_value(stdout, 0); + if (match_byte (0x31)) { - printf("footnote %d:\n", i); - dump_footnote_value(0); - match_byte(0); - match_byte(0); - match_byte(0); - match_byte(0); - if (match_byte (1)) - { - unsigned int format; - double value; - - if (match_byte (0x31)) - { - if (match_u32 (1)) - { - printf("(footnote %d) ", get_u32()); - match_byte_assert (0); - match_byte_assert (0); - int subn = get_u32 (); - printf ("nested %d bytes", subn); - pos += subn; - } - } - else - match_byte_assert (0x58); - format = get_u32 (); - value = get_double (); - printf ("value %g format %d(%d.%d)", value, format >> 16, (format >> 8) & 0xff, format & 0xff); - match_byte (1); - match_byte (0); - match_byte (0); - match_byte (0); - match_byte (1); - } - else if (match_byte (0x31)) - { - match_byte_assert(3); - get_string(); - match_byte_assert(0x58); - match_u32_assert(0); - get_string(); - match_byte(0); - } - else - match_byte_assert (0x58); - printf("(%d)\n", get_u32()); + /* Custom footnote marker string. */ + match_byte_assert(3); + get_string(); + match_byte_assert(0x58); + match_u32_assert(0); + get_string(); } + else + match_byte_assert (0x58); + printf("(%d)\n", get_u32()); + printf ("\n"); } } -static int -find_dimensions(void) -{ - { - const char dimensions[] = "-,,,.\0"; - int x = try_find_tail(dimensions, sizeof dimensions - 1); - if (x) - return x; - } - - const char dimensions[] = "-,,, .\0"; - return find_tail(dimensions, sizeof dimensions - 1); -} - static void dump_fonts(void) { - printf("fonts: offset=%08x\n", pos); match_byte(0); for (int i = 1; i <= 8; i++) { - printf("%08x: font %d, ", pos, i); + printf ("