forked from explosion/spaCy
-
Notifications
You must be signed in to change notification settings - Fork 1
/
conftest.py
193 lines (121 loc) · 4.7 KB
/
conftest.py
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
79
80
81
82
83
84
85
86
87
88
89
90
91
92
93
94
95
96
97
98
99
100
101
102
103
104
105
106
107
108
109
110
111
112
113
114
115
116
117
118
119
120
121
122
123
124
125
126
127
128
129
130
131
132
133
134
135
136
137
138
139
140
141
142
143
144
145
146
147
148
149
150
151
152
153
154
155
156
157
158
159
160
161
162
163
164
165
166
167
168
169
170
171
172
173
174
175
176
177
178
179
180
181
182
183
184
185
186
187
188
189
190
191
192
193
# coding: utf-8
from __future__ import unicode_literals
import pytest
from spacy.util import get_lang_class
def pytest_addoption(parser):
parser.addoption("--slow", action="store_true", help="include slow tests")
def pytest_runtest_setup(item):
def getopt(opt):
# When using 'pytest --pyargs spacy' to test an installed copy of
# spacy, pytest skips running our pytest_addoption() hook. Later, when
# we call getoption(), pytest raises an error, because it doesn't
# recognize the option we're asking about. To avoid this, we need to
# pass a default value. We default to False, i.e., we act like all the
# options weren't given.
return item.config.getoption("--%s" % opt, False)
for opt in ["slow"]:
if opt in item.keywords and not getopt(opt):
pytest.skip("need --%s option to run" % opt)
# Fixtures for language tokenizers (languages sorted alphabetically)
@pytest.fixture(scope="module")
def tokenizer():
return get_lang_class("xx").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def ar_tokenizer():
return get_lang_class("ar").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def bn_tokenizer():
return get_lang_class("bn").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def ca_tokenizer():
return get_lang_class("ca").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def da_tokenizer():
return get_lang_class("da").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def de_tokenizer():
return get_lang_class("de").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def el_tokenizer():
return get_lang_class("el").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def en_tokenizer():
return get_lang_class("en").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def en_vocab():
return get_lang_class("en").Defaults.create_vocab()
@pytest.fixture(scope="session")
def en_parser(en_vocab):
nlp = get_lang_class("en")(en_vocab)
return nlp.create_pipe("parser")
@pytest.fixture(scope="session")
def es_tokenizer():
return get_lang_class("es").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def fi_tokenizer():
return get_lang_class("fi").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def fr_tokenizer():
return get_lang_class("fr").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def ga_tokenizer():
return get_lang_class("ga").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def he_tokenizer():
return get_lang_class("he").Defaults.create_tokenizer()
@pytest.fixture
def hu_tokenizer():
return get_lang_class("hu").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def id_tokenizer():
return get_lang_class("id").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def it_tokenizer():
return get_lang_class("it").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def ja_tokenizer():
pytest.importorskip("MeCab")
return get_lang_class("ja").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def nb_tokenizer():
return get_lang_class("nb").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def nl_tokenizer():
return get_lang_class("nl").Defaults.create_tokenizer()
@pytest.fixture
def nl_lemmatizer(scope="session"):
return get_lang_class("nl").Defaults.create_lemmatizer()
@pytest.fixture(scope="session")
def pl_tokenizer():
return get_lang_class("pl").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def pt_tokenizer():
return get_lang_class("pt").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def ro_tokenizer():
return get_lang_class("ro").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def ru_tokenizer():
pytest.importorskip("pymorphy2")
return get_lang_class("ru").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def sv_tokenizer():
return get_lang_class("sv").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def th_tokenizer():
pytest.importorskip("pythainlp")
return get_lang_class("th").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def tr_tokenizer():
return get_lang_class("tr").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def tt_tokenizer():
return get_lang_class("tt").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def uk_tokenizer():
pytest.importorskip("pymorphy2")
pytest.importorskip("pymorphy2.lang")
return get_lang_class("uk").Defaults.create_tokenizer()
@pytest.fixture(scope="session")
def ur_tokenizer():
return get_lang_class("ur").Defaults.create_tokenizer()