From: Bruno Haible Date: Sun, 8 Feb 2009 19:54:44 +0000 (+0100) Subject: New module 'uniwbrk/table'. X-Git-Url: https://pintos-os.org/cgi-bin/gitweb.cgi?a=commitdiff_plain;h=38db1b35bf51a7a738f27d2660a9fdd3b46447be;p=pspp New module 'uniwbrk/table'. --- diff --git a/ChangeLog b/ChangeLog index 7a5eac93b9..ea08b916e2 100644 --- a/ChangeLog +++ b/ChangeLog @@ -1,5 +1,10 @@ 2009-02-08 Bruno Haible + New module 'uniwbrk/table'. + * modules/uniwbrk/table: New file. + * lib/uniwbrk/wbrktable.h: New file. + * lib/uniwbrk/wbrktable.c: New file. + New module 'uniwbrk/wordbreak-property'. * modules/uniwbrk/wordbreak-property: New file. * lib/uniwbrk/wordbreak-property.c: New file. diff --git a/lib/uniwbrk/wbrktable.c b/lib/uniwbrk/wbrktable.c new file mode 100644 index 0000000000..81a2323e0d --- /dev/null +++ b/lib/uniwbrk/wbrktable.c @@ -0,0 +1,52 @@ +/* Word break auxiliary table. + Copyright (C) 2009 Free Software Foundation, Inc. + Written by Bruno Haible , 2009. + + This program is free software: you can redistribute it and/or modify it + under the terms of the GNU Lesser General Public License as published + by the Free Software Foundation; either version 3 of the License, or + (at your option) any later version. + + This program is distributed in the hope that it will be useful, + but WITHOUT ANY WARRANTY; without even the implied warranty of + MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the GNU + Lesser General Public License for more details. + + You should have received a copy of the GNU Lesser General Public License + along with this program. If not, see . */ + +#include + +/* Specification. */ +#include "wbrktable.h" + +/* This table contains the following rules (see UAX #29): + + last current + + ALetter × ALetter (WB5) + ALetter × Numeric (WB9) + Numeric × ALetter (WB10) + Numeric × Numeric (WB8) + Katakana × Katakana (WB13) + (ALetter | Numeric | Katakana) × ExtendNumLet (WB13a) + ExtendNumLet × ExtendNumLet (WB13a) + ExtendNumLet × (ALetter | Numeric | Katakana) (WB13b) + */ + +const unsigned char uniwbrk_table[10][8] = +{ /* current: OTHER MIDNUMLET NUMERIC */ + /* KATAKANA MIDLETTER EXTENDNUMLET */ + /* ALETTER MIDNUM */ + /* last */ + /* WBP_OTHER */ { 1, 1, 1, 1, 1, 1, 1, 1 }, + /* WBP_KATAKANA */ { 1, 0, 1, 1, 1, 1, 1, 0 }, + /* WBP_ALETTER */ { 1, 1, 0, 1, 1, 1, 0, 0 }, + /* WBP_MIDNUMLET */ { 1, 1, 1, 1, 1, 1, 1, 1 }, + /* WBP_MIDLETTER */ { 1, 1, 1, 1, 1, 1, 1, 1 }, + /* WBP_MIDNUM */ { 1, 1, 1, 1, 1, 1, 1, 1 }, + /* WBP_NUMERIC */ { 1, 1, 0, 1, 1, 1, 0, 0 }, + /* WBP_EXTENDNUMLET */ { 1, 0, 0, 1, 1, 1, 0, 0 }, + /* WBP_EXTEND */ { 1, 1, 1, 1, 1, 1, 1, 1 }, + /* WBP_FORMAT */ { 1, 1, 1, 1, 1, 1, 1, 1 } +}; diff --git a/lib/uniwbrk/wbrktable.h b/lib/uniwbrk/wbrktable.h new file mode 100644 index 0000000000..14efee90a5 --- /dev/null +++ b/lib/uniwbrk/wbrktable.h @@ -0,0 +1,18 @@ +/* Word break auxiliary table. + Copyright (C) 2009 Free Software Foundation, Inc. + Written by Bruno Haible , 2009. + + This program is free software: you can redistribute it and/or modify it + under the terms of the GNU Lesser General Public License as published + by the Free Software Foundation; either version 3 of the License, or + (at your option) any later version. + + This program is distributed in the hope that it will be useful, + but WITHOUT ANY WARRANTY; without even the implied warranty of + MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the GNU + Lesser General Public License for more details. + + You should have received a copy of the GNU Lesser General Public License + along with this program. If not, see . */ + +extern const unsigned char uniwbrk_table[10][8]; diff --git a/modules/uniwbrk/table b/modules/uniwbrk/table new file mode 100644 index 0000000000..99392dbebb --- /dev/null +++ b/modules/uniwbrk/table @@ -0,0 +1,22 @@ +Description: + +Files: +lib/uniwbrk/wbrktable.h +lib/uniwbrk/wbrktable.c + +Depends-on: + +configure.ac: + +Makefile.am: +lib_SOURCES += uniwbrk/wbrktable.c + +Include: +"uniwbrk/wbrktable.h" + +License: +LGPL + +Maintainer: +Bruno Haible +