spaCy/spacy/tests/regression/test_issue351.py

12 lines
232 B
Python

# coding: utf-8
from __future__ import unicode_literals
import pytest
def test_issue351(en_tokenizer):
doc = en_tokenizer(" This is a cat.")
assert doc[0].idx == 0
assert len(doc[0]) == 3
assert doc[1].idx == 3