--- cdli/cdliSplitter.py 2007/12/13 19:20:45 1.7.2.9 +++ cdli/cdliSplitter.py 2008/01/02 15:52:01 1.7.2.10 @@ -29,11 +29,11 @@ komma_exceptionex=re.compile(komma_excep # grapheme boundaries #graphemeBounds="\{|\}|<|>|\(|\)|-|_|\#|,|\||\]|\[|\!|\?" graphemeBounds="\{|\}|<|>|-|_|\#|,|\]|\[|\!|\?|\"" -graphemeIgnore="<|>|\#|\||\]|\[|\!|\?" +graphemeIgnore="<|>|\#|\||\]|\[|\!|\?\*" # for words #wordBounds="<|>|\(|\)|_|\#|,|\||\]|\[|\!|\?" wordBounds="_|,|\"" -wordIgnore="<|>|\#|\||\]|\[|\!|\?" +wordIgnore="<|>|\#|\||\]|\[|\!|\?\*" class cdliSplitter: """base class for splitter.