Files
INTUIA/Programa final/spacy/lang/lij/punctuation.py
T

12 lines
269 B
Python
Raw Normal View History

2026-03-15 13:27:50 +00:00
from ..char_classes import ALPHA
from ..punctuation import TOKENIZER_INFIXES
ELISION = " ' ".strip().replace(" ", "").replace("\n", "")
_infixes = TOKENIZER_INFIXES + [
r"(?<=[{a}][{el}])(?=[{a}])".format(a=ALPHA, el=ELISION)
]
TOKENIZER_INFIXES = _infixes