Commit ·
544c693
1
Parent(s): a6b8354
Update spaCy pipeline
Browse files- README.md +26 -26
- accuracy.json +158 -158
- attribute_ruler/patterns +0 -0
- config.cfg +5 -4
- lemmatizer/model +1 -1
- meta.json +161 -161
- morphologizer/model +1 -1
- ner/model +2 -2
- parser/model +1 -1
- pt_core_news_md-any-py3-none-any.whl +2 -2
- senter/model +1 -1
- tok2vec/model +2 -2
- tokenizer +2 -1
- vocab/strings.json +2 -2
README.md
CHANGED
|
@@ -14,62 +14,62 @@ model-index:
|
|
| 14 |
metrics:
|
| 15 |
- name: NER Precision
|
| 16 |
type: precision
|
| 17 |
-
value: 0.
|
| 18 |
- name: NER Recall
|
| 19 |
type: recall
|
| 20 |
-
value: 0.
|
| 21 |
- name: NER F Score
|
| 22 |
type: f_score
|
| 23 |
-
value: 0.
|
| 24 |
- task:
|
| 25 |
name: TAG
|
| 26 |
type: token-classification
|
| 27 |
metrics:
|
| 28 |
- name: TAG (XPOS) Accuracy
|
| 29 |
type: accuracy
|
| 30 |
-
value: 0.
|
| 31 |
- task:
|
| 32 |
name: POS
|
| 33 |
type: token-classification
|
| 34 |
metrics:
|
| 35 |
- name: POS (UPOS) Accuracy
|
| 36 |
type: accuracy
|
| 37 |
-
value: 0.
|
| 38 |
- task:
|
| 39 |
name: MORPH
|
| 40 |
type: token-classification
|
| 41 |
metrics:
|
| 42 |
- name: Morph (UFeats) Accuracy
|
| 43 |
type: accuracy
|
| 44 |
-
value: 0.
|
| 45 |
- task:
|
| 46 |
name: LEMMA
|
| 47 |
type: token-classification
|
| 48 |
metrics:
|
| 49 |
- name: Lemma Accuracy
|
| 50 |
type: accuracy
|
| 51 |
-
value: 0.
|
| 52 |
- task:
|
| 53 |
name: UNLABELED_DEPENDENCIES
|
| 54 |
type: token-classification
|
| 55 |
metrics:
|
| 56 |
- name: Unlabeled Attachment Score (UAS)
|
| 57 |
type: f_score
|
| 58 |
-
value: 0.
|
| 59 |
- task:
|
| 60 |
name: LABELED_DEPENDENCIES
|
| 61 |
type: token-classification
|
| 62 |
metrics:
|
| 63 |
- name: Labeled Attachment Score (LAS)
|
| 64 |
type: f_score
|
| 65 |
-
value: 0.
|
| 66 |
- task:
|
| 67 |
name: SENTS
|
| 68 |
type: token-classification
|
| 69 |
metrics:
|
| 70 |
- name: Sentences F-Score
|
| 71 |
type: f_score
|
| 72 |
-
value: 0.
|
| 73 |
---
|
| 74 |
### Details: https://spacy.io/models/pt#pt_core_news_md
|
| 75 |
|
|
@@ -78,8 +78,8 @@ Portuguese pipeline optimized for CPU. Components: tok2vec, morphologizer, parse
|
|
| 78 |
| Feature | Description |
|
| 79 |
| --- | --- |
|
| 80 |
| **Name** | `pt_core_news_md` |
|
| 81 |
-
| **Version** | `3.
|
| 82 |
-
| **spaCy** | `>=3.
|
| 83 |
| **Default Pipeline** | `tok2vec`, `morphologizer`, `parser`, `lemmatizer`, `attribute_ruler`, `ner` |
|
| 84 |
| **Components** | `tok2vec`, `morphologizer`, `parser`, `lemmatizer`, `senter`, `attribute_ruler`, `ner` |
|
| 85 |
| **Vectors** | 500000 keys, 20000 unique vectors (300 dimensions) |
|
|
@@ -109,18 +109,18 @@ Portuguese pipeline optimized for CPU. Components: tok2vec, morphologizer, parse
|
|
| 109 |
| `TOKEN_P` | 99.88 |
|
| 110 |
| `TOKEN_R` | 99.95 |
|
| 111 |
| `TOKEN_F` | 99.92 |
|
| 112 |
-
| `POS_ACC` | 97.
|
| 113 |
-
| `MORPH_ACC` | 95.
|
| 114 |
-
| `MORPH_MICRO_P` | 97.
|
| 115 |
-
| `MORPH_MICRO_R` | 97.
|
| 116 |
-
| `MORPH_MICRO_F` | 97.
|
| 117 |
-
| `SENTS_P` |
|
| 118 |
-
| `SENTS_R` | 95.
|
| 119 |
-
| `SENTS_F` | 94.
|
| 120 |
-
| `DEP_UAS` | 90.
|
| 121 |
-
| `DEP_LAS` |
|
| 122 |
-
| `LEMMA_ACC` | 97.
|
| 123 |
| `TAG_ACC` | 89.62 |
|
| 124 |
-
| `ENTS_P` | 89.
|
| 125 |
-
| `ENTS_R` | 89.
|
| 126 |
-
| `ENTS_F` | 89.
|
|
|
|
| 14 |
metrics:
|
| 15 |
- name: NER Precision
|
| 16 |
type: precision
|
| 17 |
+
value: 0.8937657035
|
| 18 |
- name: NER Recall
|
| 19 |
type: recall
|
| 20 |
+
value: 0.8960699034
|
| 21 |
- name: NER F Score
|
| 22 |
type: f_score
|
| 23 |
+
value: 0.8949163203
|
| 24 |
- task:
|
| 25 |
name: TAG
|
| 26 |
type: token-classification
|
| 27 |
metrics:
|
| 28 |
- name: TAG (XPOS) Accuracy
|
| 29 |
type: accuracy
|
| 30 |
+
value: 0.8962306206
|
| 31 |
- task:
|
| 32 |
name: POS
|
| 33 |
type: token-classification
|
| 34 |
metrics:
|
| 35 |
- name: POS (UPOS) Accuracy
|
| 36 |
type: accuracy
|
| 37 |
+
value: 0.9702313141
|
| 38 |
- task:
|
| 39 |
name: MORPH
|
| 40 |
type: token-classification
|
| 41 |
metrics:
|
| 42 |
- name: Morph (UFeats) Accuracy
|
| 43 |
type: accuracy
|
| 44 |
+
value: 0.9534422982
|
| 45 |
- task:
|
| 46 |
name: LEMMA
|
| 47 |
type: token-classification
|
| 48 |
metrics:
|
| 49 |
- name: Lemma Accuracy
|
| 50 |
type: accuracy
|
| 51 |
+
value: 0.9703333168
|
| 52 |
- task:
|
| 53 |
name: UNLABELED_DEPENDENCIES
|
| 54 |
type: token-classification
|
| 55 |
metrics:
|
| 56 |
- name: Unlabeled Attachment Score (UAS)
|
| 57 |
type: f_score
|
| 58 |
+
value: 0.9013711257
|
| 59 |
- task:
|
| 60 |
name: LABELED_DEPENDENCIES
|
| 61 |
type: token-classification
|
| 62 |
metrics:
|
| 63 |
- name: Labeled Attachment Score (LAS)
|
| 64 |
type: f_score
|
| 65 |
+
value: 0.8593155894
|
| 66 |
- task:
|
| 67 |
name: SENTS
|
| 68 |
type: token-classification
|
| 69 |
metrics:
|
| 70 |
- name: Sentences F-Score
|
| 71 |
type: f_score
|
| 72 |
+
value: 0.9457981824
|
| 73 |
---
|
| 74 |
### Details: https://spacy.io/models/pt#pt_core_news_md
|
| 75 |
|
|
|
|
| 78 |
| Feature | Description |
|
| 79 |
| --- | --- |
|
| 80 |
| **Name** | `pt_core_news_md` |
|
| 81 |
+
| **Version** | `3.5.0` |
|
| 82 |
+
| **spaCy** | `>=3.5.0,<3.6.0` |
|
| 83 |
| **Default Pipeline** | `tok2vec`, `morphologizer`, `parser`, `lemmatizer`, `attribute_ruler`, `ner` |
|
| 84 |
| **Components** | `tok2vec`, `morphologizer`, `parser`, `lemmatizer`, `senter`, `attribute_ruler`, `ner` |
|
| 85 |
| **Vectors** | 500000 keys, 20000 unique vectors (300 dimensions) |
|
|
|
|
| 109 |
| `TOKEN_P` | 99.88 |
|
| 110 |
| `TOKEN_R` | 99.95 |
|
| 111 |
| `TOKEN_F` | 99.92 |
|
| 112 |
+
| `POS_ACC` | 97.02 |
|
| 113 |
+
| `MORPH_ACC` | 95.34 |
|
| 114 |
+
| `MORPH_MICRO_P` | 97.84 |
|
| 115 |
+
| `MORPH_MICRO_R` | 97.51 |
|
| 116 |
+
| `MORPH_MICRO_F` | 97.68 |
|
| 117 |
+
| `SENTS_P` | 93.47 |
|
| 118 |
+
| `SENTS_R` | 95.98 |
|
| 119 |
+
| `SENTS_F` | 94.58 |
|
| 120 |
+
| `DEP_UAS` | 90.14 |
|
| 121 |
+
| `DEP_LAS` | 85.93 |
|
| 122 |
+
| `LEMMA_ACC` | 97.03 |
|
| 123 |
| `TAG_ACC` | 89.62 |
|
| 124 |
+
| `ENTS_P` | 89.38 |
|
| 125 |
+
| `ENTS_R` | 89.61 |
|
| 126 |
+
| `ENTS_F` | 89.49 |
|
accuracy.json
CHANGED
|
@@ -3,148 +3,148 @@
|
|
| 3 |
"token_p": 0.9988117635,
|
| 4 |
"token_r": 0.9995045581,
|
| 5 |
"token_f": 0.9991580407,
|
| 6 |
-
"pos_acc": 0.
|
| 7 |
-
"morph_acc": 0.
|
| 8 |
-
"morph_micro_p": 0.
|
| 9 |
-
"morph_micro_r": 0.
|
| 10 |
-
"morph_micro_f": 0.
|
| 11 |
"morph_per_feat": {
|
| 12 |
"Mood": {
|
| 13 |
-
"p": 0.
|
| 14 |
"r": 0.9794437727,
|
| 15 |
-
"f": 0.
|
| 16 |
},
|
| 17 |
"Number": {
|
| 18 |
-
"p": 0.
|
| 19 |
-
"r": 0.
|
| 20 |
-
"f": 0.
|
| 21 |
},
|
| 22 |
"Person": {
|
| 23 |
-
"p": 0.
|
| 24 |
-
"r": 0.
|
| 25 |
-
"f": 0.
|
| 26 |
},
|
| 27 |
"Tense": {
|
| 28 |
-
"p": 0.
|
| 29 |
-
"r": 0.
|
| 30 |
-
"f": 0.
|
| 31 |
},
|
| 32 |
"VerbForm": {
|
| 33 |
-
"p": 0.
|
| 34 |
-
"r": 0.
|
| 35 |
-
"f": 0.
|
| 36 |
},
|
| 37 |
"Gender": {
|
| 38 |
-
"p": 0.
|
| 39 |
-
"r": 0.
|
| 40 |
-
"f": 0.
|
| 41 |
},
|
| 42 |
"PronType": {
|
| 43 |
-
"p": 0.
|
| 44 |
-
"r": 0.
|
| 45 |
-
"f": 0.
|
| 46 |
},
|
| 47 |
"Definite": {
|
| 48 |
-
"p": 0.
|
| 49 |
-
"r": 0.
|
| 50 |
-
"f": 0.
|
| 51 |
},
|
| 52 |
"NumType": {
|
| 53 |
-
"p": 0.
|
| 54 |
-
"r": 0.
|
| 55 |
-
"f": 0.
|
| 56 |
},
|
| 57 |
"Voice": {
|
| 58 |
-
"p": 0.
|
| 59 |
"r": 0.9302325581,
|
| 60 |
-
"f": 0.
|
| 61 |
},
|
| 62 |
"Polarity": {
|
| 63 |
-
"p":
|
| 64 |
-
"r":
|
| 65 |
-
"f": 0.
|
| 66 |
},
|
| 67 |
"Case": {
|
| 68 |
-
"p": 0.
|
| 69 |
-
"r": 0.
|
| 70 |
-
"f": 0.
|
| 71 |
}
|
| 72 |
},
|
| 73 |
-
"sents_p": 0.
|
| 74 |
-
"sents_r": 0.
|
| 75 |
-
"sents_f": 0.
|
| 76 |
-
"dep_uas": 0.
|
| 77 |
-
"dep_las": 0.
|
| 78 |
"dep_las_per_type": {
|
| 79 |
"cop": {
|
| 80 |
-
"p": 0.
|
| 81 |
"r": 0.9361702128,
|
| 82 |
-
"f": 0.
|
| 83 |
},
|
| 84 |
"root": {
|
| 85 |
-
"p": 0.
|
| 86 |
-
"r": 0.
|
| 87 |
-
"f": 0.
|
| 88 |
},
|
| 89 |
"det": {
|
| 90 |
-
"p": 0.
|
| 91 |
-
"r": 0.
|
| 92 |
-
"f": 0.
|
| 93 |
},
|
| 94 |
"amod": {
|
| 95 |
-
"p": 0.
|
| 96 |
-
"r": 0.
|
| 97 |
-
"f": 0.
|
| 98 |
},
|
| 99 |
"nsubj": {
|
| 100 |
-
"p": 0.
|
| 101 |
-
"r": 0.
|
| 102 |
-
"f": 0.
|
| 103 |
},
|
| 104 |
"case": {
|
| 105 |
-
"p": 0.
|
| 106 |
-
"r": 0.
|
| 107 |
-
"f": 0.
|
| 108 |
},
|
| 109 |
"nmod": {
|
| 110 |
-
"p": 0.
|
| 111 |
-
"r": 0.
|
| 112 |
-
"f": 0.
|
| 113 |
},
|
| 114 |
"flat:name": {
|
| 115 |
-
"p": 0.
|
| 116 |
-
"r": 0.
|
| 117 |
-
"f": 0.
|
| 118 |
},
|
| 119 |
"acl": {
|
| 120 |
-
"p": 0.
|
| 121 |
-
"r": 0.
|
| 122 |
-
"f": 0.
|
| 123 |
},
|
| 124 |
"advmod": {
|
| 125 |
-
"p": 0.
|
| 126 |
-
"r": 0.
|
| 127 |
-
"f": 0.
|
| 128 |
},
|
| 129 |
"obl": {
|
| 130 |
-
"p": 0.
|
| 131 |
-
"r": 0.
|
| 132 |
-
"f": 0.
|
| 133 |
},
|
| 134 |
"conj": {
|
| 135 |
-
"p": 0.
|
| 136 |
-
"r": 0.
|
| 137 |
-
"f": 0.
|
| 138 |
},
|
| 139 |
"cc": {
|
| 140 |
-
"p": 0.
|
| 141 |
-
"r": 0.
|
| 142 |
-
"f": 0.
|
| 143 |
},
|
| 144 |
"obj": {
|
| 145 |
-
"p": 0.
|
| 146 |
-
"r": 0.
|
| 147 |
-
"f": 0.
|
| 148 |
},
|
| 149 |
"obl:agent": {
|
| 150 |
"p": 0.7941176471,
|
|
@@ -152,59 +152,64 @@
|
|
| 152 |
"f": 0.84375
|
| 153 |
},
|
| 154 |
"acl:relcl": {
|
| 155 |
-
"p": 0.
|
| 156 |
-
"r": 0.
|
| 157 |
-
"f": 0.
|
| 158 |
},
|
| 159 |
"mark": {
|
| 160 |
-
"p": 0.
|
| 161 |
-
"r": 0.
|
| 162 |
-
"f": 0.
|
| 163 |
},
|
| 164 |
"advcl": {
|
| 165 |
-
"p": 0.
|
| 166 |
-
"r": 0.
|
| 167 |
-
"f": 0.
|
| 168 |
},
|
| 169 |
"xcomp": {
|
| 170 |
-
"p": 0.
|
| 171 |
-
"r": 0.
|
| 172 |
-
"f": 0.
|
| 173 |
},
|
| 174 |
"iobj": {
|
| 175 |
-
"p": 0.
|
| 176 |
-
"r": 0.
|
| 177 |
-
"f": 0.
|
| 178 |
},
|
| 179 |
"appos": {
|
| 180 |
-
"p": 0.
|
| 181 |
-
"r": 0.
|
| 182 |
-
"f": 0.
|
| 183 |
},
|
| 184 |
"fixed": {
|
| 185 |
-
"p": 0.
|
| 186 |
-
"r": 0.
|
| 187 |
-
"f": 0.
|
| 188 |
},
|
| 189 |
"nummod": {
|
| 190 |
-
"p": 0.
|
| 191 |
-
"r": 0.
|
| 192 |
-
"f": 0.
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 193 |
},
|
| 194 |
"aux": {
|
| 195 |
-
"p": 0.
|
| 196 |
-
"r": 0.
|
| 197 |
-
"f": 0.
|
| 198 |
},
|
| 199 |
"csubj": {
|
| 200 |
-
"p": 0.
|
| 201 |
-
"r": 0.
|
| 202 |
-
"f": 0.
|
| 203 |
},
|
| 204 |
"ccomp": {
|
| 205 |
-
"p": 0.
|
| 206 |
-
"r": 0.
|
| 207 |
-
"f": 0.
|
| 208 |
},
|
| 209 |
"orphan": {
|
| 210 |
"p": 0.0,
|
|
@@ -218,18 +223,18 @@
|
|
| 218 |
},
|
| 219 |
"aux:pass": {
|
| 220 |
"p": 1.0,
|
| 221 |
-
"r": 0.
|
| 222 |
-
"f": 0.
|
| 223 |
},
|
| 224 |
"nsubj:pass": {
|
| 225 |
-
"p": 0.
|
| 226 |
-
"r": 0.
|
| 227 |
-
"f": 0.
|
| 228 |
},
|
| 229 |
"parataxis": {
|
| 230 |
-
"p": 0.
|
| 231 |
-
"r": 0.
|
| 232 |
-
"f": 0.
|
| 233 |
},
|
| 234 |
"list": {
|
| 235 |
"p": 0.0,
|
|
@@ -237,19 +242,14 @@
|
|
| 237 |
"f": 0.0
|
| 238 |
},
|
| 239 |
"expl": {
|
| 240 |
-
"p": 0.
|
| 241 |
-
"r":
|
| 242 |
-
"f": 0.
|
| 243 |
},
|
| 244 |
"compound": {
|
| 245 |
-
"p": 0.
|
| 246 |
-
"r": 0.
|
| 247 |
-
"f": 0.
|
| 248 |
-
},
|
| 249 |
-
"dep": {
|
| 250 |
-
"p": 0.0,
|
| 251 |
-
"r": 0.0,
|
| 252 |
-
"f": 0.0
|
| 253 |
},
|
| 254 |
"vocative": {
|
| 255 |
"p": 0.0,
|
|
@@ -257,9 +257,9 @@
|
|
| 257 |
"f": 0.0
|
| 258 |
},
|
| 259 |
"discourse": {
|
| 260 |
-
"p":
|
| 261 |
-
"r": 0.
|
| 262 |
-
"f": 0.
|
| 263 |
},
|
| 264 |
"expl:pass": {
|
| 265 |
"p": 0.0,
|
|
@@ -272,32 +272,32 @@
|
|
| 272 |
"f": 0.8571428571
|
| 273 |
}
|
| 274 |
},
|
| 275 |
-
"lemma_acc": 0.
|
| 276 |
-
"tag_acc": 0.
|
| 277 |
-
"ents_p": 0.
|
| 278 |
-
"ents_r": 0.
|
| 279 |
-
"ents_f": 0.
|
| 280 |
"ents_per_type": {
|
| 281 |
"LOC": {
|
| 282 |
-
"p": 0.
|
| 283 |
-
"r": 0.
|
| 284 |
-
"f": 0.
|
| 285 |
},
|
| 286 |
"PER": {
|
| 287 |
-
"p": 0.
|
| 288 |
-
"r": 0.
|
| 289 |
-
"f": 0.
|
| 290 |
},
|
| 291 |
"ORG": {
|
| 292 |
-
"p": 0.
|
| 293 |
-
"r": 0.
|
| 294 |
-
"f": 0.
|
| 295 |
},
|
| 296 |
"MISC": {
|
| 297 |
-
"p": 0.
|
| 298 |
-
"r": 0.
|
| 299 |
-
"f": 0.
|
| 300 |
}
|
| 301 |
},
|
| 302 |
-
"speed":
|
| 303 |
}
|
|
|
|
| 3 |
"token_p": 0.9988117635,
|
| 4 |
"token_r": 0.9995045581,
|
| 5 |
"token_f": 0.9991580407,
|
| 6 |
+
"pos_acc": 0.9702313141,
|
| 7 |
+
"morph_acc": 0.9534422982,
|
| 8 |
+
"morph_micro_p": 0.978440699,
|
| 9 |
+
"morph_micro_r": 0.9751133553,
|
| 10 |
+
"morph_micro_f": 0.9767741935,
|
| 11 |
"morph_per_feat": {
|
| 12 |
"Mood": {
|
| 13 |
+
"p": 0.9770808203,
|
| 14 |
"r": 0.9794437727,
|
| 15 |
+
"f": 0.9782608696
|
| 16 |
},
|
| 17 |
"Number": {
|
| 18 |
+
"p": 0.9926424546,
|
| 19 |
+
"r": 0.9875408815,
|
| 20 |
+
"f": 0.9900850964
|
| 21 |
},
|
| 22 |
"Person": {
|
| 23 |
+
"p": 0.9712765957,
|
| 24 |
+
"r": 0.9743863394,
|
| 25 |
+
"f": 0.9728289824
|
| 26 |
},
|
| 27 |
"Tense": {
|
| 28 |
+
"p": 0.9304677623,
|
| 29 |
+
"r": 0.9684210526,
|
| 30 |
+
"f": 0.9490651193
|
| 31 |
},
|
| 32 |
"VerbForm": {
|
| 33 |
+
"p": 0.9732072498,
|
| 34 |
+
"r": 0.9786053883,
|
| 35 |
+
"f": 0.9758988542
|
| 36 |
},
|
| 37 |
"Gender": {
|
| 38 |
+
"p": 0.9659508609,
|
| 39 |
+
"r": 0.9548670874,
|
| 40 |
+
"f": 0.9603769956
|
| 41 |
},
|
| 42 |
"PronType": {
|
| 43 |
+
"p": 0.9890625,
|
| 44 |
+
"r": 0.9813953488,
|
| 45 |
+
"f": 0.9852140078
|
| 46 |
},
|
| 47 |
"Definite": {
|
| 48 |
+
"p": 0.9952574526,
|
| 49 |
+
"r": 0.9952574526,
|
| 50 |
+
"f": 0.9952574526
|
| 51 |
},
|
| 52 |
"NumType": {
|
| 53 |
+
"p": 0.9671052632,
|
| 54 |
+
"r": 0.9639344262,
|
| 55 |
+
"f": 0.9655172414
|
| 56 |
},
|
| 57 |
"Voice": {
|
| 58 |
+
"p": 0.9195402299,
|
| 59 |
"r": 0.9302325581,
|
| 60 |
+
"f": 0.9248554913
|
| 61 |
},
|
| 62 |
"Polarity": {
|
| 63 |
+
"p": 1.0,
|
| 64 |
+
"r": 0.9861111111,
|
| 65 |
+
"f": 0.993006993
|
| 66 |
},
|
| 67 |
"Case": {
|
| 68 |
+
"p": 0.8571428571,
|
| 69 |
+
"r": 0.8571428571,
|
| 70 |
+
"f": 0.8571428571
|
| 71 |
}
|
| 72 |
},
|
| 73 |
+
"sents_p": 0.9346642468,
|
| 74 |
+
"sents_r": 0.9597677922,
|
| 75 |
+
"sents_f": 0.9457981824,
|
| 76 |
+
"dep_uas": 0.9013711257,
|
| 77 |
+
"dep_las": 0.8593155894,
|
| 78 |
"dep_las_per_type": {
|
| 79 |
"cop": {
|
| 80 |
+
"p": 0.8684210526,
|
| 81 |
"r": 0.9361702128,
|
| 82 |
+
"f": 0.9010238908
|
| 83 |
},
|
| 84 |
"root": {
|
| 85 |
+
"p": 0.9364791289,
|
| 86 |
+
"r": 0.9214285714,
|
| 87 |
+
"f": 0.9288928893
|
| 88 |
},
|
| 89 |
"det": {
|
| 90 |
+
"p": 0.9638554217,
|
| 91 |
+
"r": 0.9659714599,
|
| 92 |
+
"f": 0.9649122807
|
| 93 |
},
|
| 94 |
"amod": {
|
| 95 |
+
"p": 0.9423558897,
|
| 96 |
+
"r": 0.9170731707,
|
| 97 |
+
"f": 0.9295426452
|
| 98 |
},
|
| 99 |
"nsubj": {
|
| 100 |
+
"p": 0.8890814558,
|
| 101 |
+
"r": 0.8952879581,
|
| 102 |
+
"f": 0.892173913
|
| 103 |
},
|
| 104 |
"case": {
|
| 105 |
+
"p": 0.973265436,
|
| 106 |
+
"r": 0.9757498405,
|
| 107 |
+
"f": 0.9745060548
|
| 108 |
},
|
| 109 |
"nmod": {
|
| 110 |
+
"p": 0.8051948052,
|
| 111 |
+
"r": 0.8285077951,
|
| 112 |
+
"f": 0.8166849616
|
| 113 |
},
|
| 114 |
"flat:name": {
|
| 115 |
+
"p": 0.8732394366,
|
| 116 |
+
"r": 0.9018181818,
|
| 117 |
+
"f": 0.8872987478
|
| 118 |
},
|
| 119 |
"acl": {
|
| 120 |
+
"p": 0.6476190476,
|
| 121 |
+
"r": 0.5862068966,
|
| 122 |
+
"f": 0.6153846154
|
| 123 |
},
|
| 124 |
"advmod": {
|
| 125 |
+
"p": 0.8163841808,
|
| 126 |
+
"r": 0.8210227273,
|
| 127 |
+
"f": 0.8186968839
|
| 128 |
},
|
| 129 |
"obl": {
|
| 130 |
+
"p": 0.6811881188,
|
| 131 |
+
"r": 0.70781893,
|
| 132 |
+
"f": 0.6942482341
|
| 133 |
},
|
| 134 |
"conj": {
|
| 135 |
+
"p": 0.5871212121,
|
| 136 |
+
"r": 0.5719557196,
|
| 137 |
+
"f": 0.5794392523
|
| 138 |
},
|
| 139 |
"cc": {
|
| 140 |
+
"p": 0.8808510638,
|
| 141 |
+
"r": 0.9,
|
| 142 |
+
"f": 0.8903225806
|
| 143 |
},
|
| 144 |
"obj": {
|
| 145 |
+
"p": 0.8734693878,
|
| 146 |
+
"r": 0.8246628131,
|
| 147 |
+
"f": 0.8483647175
|
| 148 |
},
|
| 149 |
"obl:agent": {
|
| 150 |
"p": 0.7941176471,
|
|
|
|
| 152 |
"f": 0.84375
|
| 153 |
},
|
| 154 |
"acl:relcl": {
|
| 155 |
+
"p": 0.610619469,
|
| 156 |
+
"r": 0.6509433962,
|
| 157 |
+
"f": 0.6301369863
|
| 158 |
},
|
| 159 |
"mark": {
|
| 160 |
+
"p": 0.8883495146,
|
| 161 |
+
"r": 0.8551401869,
|
| 162 |
+
"f": 0.8714285714
|
| 163 |
},
|
| 164 |
"advcl": {
|
| 165 |
+
"p": 0.572519084,
|
| 166 |
+
"r": 0.6696428571,
|
| 167 |
+
"f": 0.6172839506
|
| 168 |
},
|
| 169 |
"xcomp": {
|
| 170 |
+
"p": 0.8482142857,
|
| 171 |
+
"r": 0.76,
|
| 172 |
+
"f": 0.8016877637
|
| 173 |
},
|
| 174 |
"iobj": {
|
| 175 |
+
"p": 0.6666666667,
|
| 176 |
+
"r": 0.1818181818,
|
| 177 |
+
"f": 0.2857142857
|
| 178 |
},
|
| 179 |
"appos": {
|
| 180 |
+
"p": 0.5895953757,
|
| 181 |
+
"r": 0.6219512195,
|
| 182 |
+
"f": 0.6053412463
|
| 183 |
},
|
| 184 |
"fixed": {
|
| 185 |
+
"p": 0.7692307692,
|
| 186 |
+
"r": 0.7142857143,
|
| 187 |
+
"f": 0.7407407407
|
| 188 |
},
|
| 189 |
"nummod": {
|
| 190 |
+
"p": 0.9671052632,
|
| 191 |
+
"r": 0.9423076923,
|
| 192 |
+
"f": 0.9545454545
|
| 193 |
+
},
|
| 194 |
+
"dep": {
|
| 195 |
+
"p": 0.0,
|
| 196 |
+
"r": 0.0,
|
| 197 |
+
"f": 0.0
|
| 198 |
},
|
| 199 |
"aux": {
|
| 200 |
+
"p": 0.9076923077,
|
| 201 |
+
"r": 0.9516129032,
|
| 202 |
+
"f": 0.9291338583
|
| 203 |
},
|
| 204 |
"csubj": {
|
| 205 |
+
"p": 0.6666666667,
|
| 206 |
+
"r": 0.5,
|
| 207 |
+
"f": 0.5714285714
|
| 208 |
},
|
| 209 |
"ccomp": {
|
| 210 |
+
"p": 0.75,
|
| 211 |
+
"r": 0.65,
|
| 212 |
+
"f": 0.6964285714
|
| 213 |
},
|
| 214 |
"orphan": {
|
| 215 |
"p": 0.0,
|
|
|
|
| 223 |
},
|
| 224 |
"aux:pass": {
|
| 225 |
"p": 1.0,
|
| 226 |
+
"r": 0.9705882353,
|
| 227 |
+
"f": 0.9850746269
|
| 228 |
},
|
| 229 |
"nsubj:pass": {
|
| 230 |
+
"p": 0.8679245283,
|
| 231 |
+
"r": 0.8518518519,
|
| 232 |
+
"f": 0.8598130841
|
| 233 |
},
|
| 234 |
"parataxis": {
|
| 235 |
+
"p": 0.3818181818,
|
| 236 |
+
"r": 0.2957746479,
|
| 237 |
+
"f": 0.3333333333
|
| 238 |
},
|
| 239 |
"list": {
|
| 240 |
"p": 0.0,
|
|
|
|
| 242 |
"f": 0.0
|
| 243 |
},
|
| 244 |
"expl": {
|
| 245 |
+
"p": 0.76,
|
| 246 |
+
"r": 1.0,
|
| 247 |
+
"f": 0.8636363636
|
| 248 |
},
|
| 249 |
"compound": {
|
| 250 |
+
"p": 0.75,
|
| 251 |
+
"r": 0.6,
|
| 252 |
+
"f": 0.6666666667
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 253 |
},
|
| 254 |
"vocative": {
|
| 255 |
"p": 0.0,
|
|
|
|
| 257 |
"f": 0.0
|
| 258 |
},
|
| 259 |
"discourse": {
|
| 260 |
+
"p": 1.0,
|
| 261 |
+
"r": 0.3333333333,
|
| 262 |
+
"f": 0.5
|
| 263 |
},
|
| 264 |
"expl:pass": {
|
| 265 |
"p": 0.0,
|
|
|
|
| 272 |
"f": 0.8571428571
|
| 273 |
}
|
| 274 |
},
|
| 275 |
+
"lemma_acc": 0.9703333168,
|
| 276 |
+
"tag_acc": 0.8962306206,
|
| 277 |
+
"ents_p": 0.8937657035,
|
| 278 |
+
"ents_r": 0.8960699034,
|
| 279 |
+
"ents_f": 0.8949163203,
|
| 280 |
"ents_per_type": {
|
| 281 |
"LOC": {
|
| 282 |
+
"p": 0.9114664452,
|
| 283 |
+
"r": 0.9259845288,
|
| 284 |
+
"f": 0.9186681318
|
| 285 |
},
|
| 286 |
"PER": {
|
| 287 |
+
"p": 0.90687251,
|
| 288 |
+
"r": 0.9220253165,
|
| 289 |
+
"f": 0.9143861411
|
| 290 |
},
|
| 291 |
"ORG": {
|
| 292 |
+
"p": 0.8461361015,
|
| 293 |
+
"r": 0.8192986375,
|
| 294 |
+
"f": 0.8325011348
|
| 295 |
},
|
| 296 |
"MISC": {
|
| 297 |
+
"p": 0.8121030346,
|
| 298 |
+
"r": 0.7613298048,
|
| 299 |
+
"f": 0.785897217
|
| 300 |
}
|
| 301 |
},
|
| 302 |
+
"speed": 11483.8429749162
|
| 303 |
}
|
attribute_ruler/patterns
CHANGED
|
Binary files a/attribute_ruler/patterns and b/attribute_ruler/patterns differ
|
|
|
config.cfg
CHANGED
|
@@ -81,8 +81,8 @@ nO = null
|
|
| 81 |
[components.ner.model.tok2vec.embed]
|
| 82 |
@architectures = "spacy.MultiHashEmbed.v2"
|
| 83 |
width = 96
|
| 84 |
-
attrs = ["NORM","PREFIX","SUFFIX","SHAPE"
|
| 85 |
-
rows = [5000,1000,2500,2500
|
| 86 |
include_static_vectors = true
|
| 87 |
|
| 88 |
[components.ner.model.tok2vec.encode]
|
|
@@ -150,8 +150,8 @@ factory = "tok2vec"
|
|
| 150 |
[components.tok2vec.model.embed]
|
| 151 |
@architectures = "spacy.MultiHashEmbed.v2"
|
| 152 |
width = ${components.tok2vec.model.encode:width}
|
| 153 |
-
attrs = ["NORM","PREFIX","SUFFIX","SHAPE","SPACY"]
|
| 154 |
-
rows = [5000,1000,2500,2500,50]
|
| 155 |
include_static_vectors = true
|
| 156 |
|
| 157 |
[components.tok2vec.model.encode]
|
|
@@ -193,6 +193,7 @@ eval_frequency = 1000
|
|
| 193 |
frozen_components = []
|
| 194 |
before_to_disk = null
|
| 195 |
annotating_components = []
|
|
|
|
| 196 |
|
| 197 |
[training.batcher]
|
| 198 |
@batchers = "spacy.batch_by_words.v1"
|
|
|
|
| 81 |
[components.ner.model.tok2vec.embed]
|
| 82 |
@architectures = "spacy.MultiHashEmbed.v2"
|
| 83 |
width = 96
|
| 84 |
+
attrs = ["NORM","PREFIX","SUFFIX","SHAPE"]
|
| 85 |
+
rows = [5000,1000,2500,2500]
|
| 86 |
include_static_vectors = true
|
| 87 |
|
| 88 |
[components.ner.model.tok2vec.encode]
|
|
|
|
| 150 |
[components.tok2vec.model.embed]
|
| 151 |
@architectures = "spacy.MultiHashEmbed.v2"
|
| 152 |
width = ${components.tok2vec.model.encode:width}
|
| 153 |
+
attrs = ["NORM","PREFIX","SUFFIX","SHAPE","SPACY","IS_SPACE"]
|
| 154 |
+
rows = [5000,1000,2500,2500,50,50]
|
| 155 |
include_static_vectors = true
|
| 156 |
|
| 157 |
[components.tok2vec.model.encode]
|
|
|
|
| 193 |
frozen_components = []
|
| 194 |
before_to_disk = null
|
| 195 |
annotating_components = []
|
| 196 |
+
before_update = null
|
| 197 |
|
| 198 |
[training.batcher]
|
| 199 |
@batchers = "spacy.batch_by_words.v1"
|
lemmatizer/model
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 217334
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:f5e9ebedbb4e398a629f6d3c2c89cbf1b016138a4173bc148422e48c16b82615
|
| 3 |
size 217334
|
meta.json
CHANGED
|
@@ -1,14 +1,14 @@
|
|
| 1 |
{
|
| 2 |
"lang":"pt",
|
| 3 |
"name":"core_news_md",
|
| 4 |
-
"version":"3.
|
| 5 |
"description":"Portuguese pipeline optimized for CPU. Components: tok2vec, morphologizer, parser, lemmatizer (trainable_lemmatizer), senter, ner, attribute_ruler.",
|
| 6 |
"author":"Explosion",
|
| 7 |
"email":"contact@explosion.ai",
|
| 8 |
"url":"https://explosion.ai",
|
| 9 |
"license":"CC BY-SA 4.0",
|
| 10 |
-
"spacy_version":">=3.
|
| 11 |
-
"spacy_git_version":"
|
| 12 |
"vectors":{
|
| 13 |
"width":300,
|
| 14 |
"vectors":20000,
|
|
@@ -644,148 +644,148 @@
|
|
| 644 |
"token_p":0.9988117635,
|
| 645 |
"token_r":0.9995045581,
|
| 646 |
"token_f":0.9991580407,
|
| 647 |
-
"pos_acc":0.
|
| 648 |
-
"morph_acc":0.
|
| 649 |
-
"morph_micro_p":0.
|
| 650 |
-
"morph_micro_r":0.
|
| 651 |
-
"morph_micro_f":0.
|
| 652 |
"morph_per_feat":{
|
| 653 |
"Mood":{
|
| 654 |
-
"p":0.
|
| 655 |
"r":0.9794437727,
|
| 656 |
-
"f":0.
|
| 657 |
},
|
| 658 |
"Number":{
|
| 659 |
-
"p":0.
|
| 660 |
-
"r":0.
|
| 661 |
-
"f":0.
|
| 662 |
},
|
| 663 |
"Person":{
|
| 664 |
-
"p":0.
|
| 665 |
-
"r":0.
|
| 666 |
-
"f":0.
|
| 667 |
},
|
| 668 |
"Tense":{
|
| 669 |
-
"p":0.
|
| 670 |
-
"r":0.
|
| 671 |
-
"f":0.
|
| 672 |
},
|
| 673 |
"VerbForm":{
|
| 674 |
-
"p":0.
|
| 675 |
-
"r":0.
|
| 676 |
-
"f":0.
|
| 677 |
},
|
| 678 |
"Gender":{
|
| 679 |
-
"p":0.
|
| 680 |
-
"r":0.
|
| 681 |
-
"f":0.
|
| 682 |
},
|
| 683 |
"PronType":{
|
| 684 |
-
"p":0.
|
| 685 |
-
"r":0.
|
| 686 |
-
"f":0.
|
| 687 |
},
|
| 688 |
"Definite":{
|
| 689 |
-
"p":0.
|
| 690 |
-
"r":0.
|
| 691 |
-
"f":0.
|
| 692 |
},
|
| 693 |
"NumType":{
|
| 694 |
-
"p":0.
|
| 695 |
-
"r":0.
|
| 696 |
-
"f":0.
|
| 697 |
},
|
| 698 |
"Voice":{
|
| 699 |
-
"p":0.
|
| 700 |
"r":0.9302325581,
|
| 701 |
-
"f":0.
|
| 702 |
},
|
| 703 |
"Polarity":{
|
| 704 |
-
"p":
|
| 705 |
-
"r":
|
| 706 |
-
"f":0.
|
| 707 |
},
|
| 708 |
"Case":{
|
| 709 |
-
"p":0.
|
| 710 |
-
"r":0.
|
| 711 |
-
"f":0.
|
| 712 |
}
|
| 713 |
},
|
| 714 |
-
"sents_p":0.
|
| 715 |
-
"sents_r":0.
|
| 716 |
-
"sents_f":0.
|
| 717 |
-
"dep_uas":0.
|
| 718 |
-
"dep_las":0.
|
| 719 |
"dep_las_per_type":{
|
| 720 |
"cop":{
|
| 721 |
-
"p":0.
|
| 722 |
"r":0.9361702128,
|
| 723 |
-
"f":0.
|
| 724 |
},
|
| 725 |
"root":{
|
| 726 |
-
"p":0.
|
| 727 |
-
"r":0.
|
| 728 |
-
"f":0.
|
| 729 |
},
|
| 730 |
"det":{
|
| 731 |
-
"p":0.
|
| 732 |
-
"r":0.
|
| 733 |
-
"f":0.
|
| 734 |
},
|
| 735 |
"amod":{
|
| 736 |
-
"p":0.
|
| 737 |
-
"r":0.
|
| 738 |
-
"f":0.
|
| 739 |
},
|
| 740 |
"nsubj":{
|
| 741 |
-
"p":0.
|
| 742 |
-
"r":0.
|
| 743 |
-
"f":0.
|
| 744 |
},
|
| 745 |
"case":{
|
| 746 |
-
"p":0.
|
| 747 |
-
"r":0.
|
| 748 |
-
"f":0.
|
| 749 |
},
|
| 750 |
"nmod":{
|
| 751 |
-
"p":0.
|
| 752 |
-
"r":0.
|
| 753 |
-
"f":0.
|
| 754 |
},
|
| 755 |
"flat:name":{
|
| 756 |
-
"p":0.
|
| 757 |
-
"r":0.
|
| 758 |
-
"f":0.
|
| 759 |
},
|
| 760 |
"acl":{
|
| 761 |
-
"p":0.
|
| 762 |
-
"r":0.
|
| 763 |
-
"f":0.
|
| 764 |
},
|
| 765 |
"advmod":{
|
| 766 |
-
"p":0.
|
| 767 |
-
"r":0.
|
| 768 |
-
"f":0.
|
| 769 |
},
|
| 770 |
"obl":{
|
| 771 |
-
"p":0.
|
| 772 |
-
"r":0.
|
| 773 |
-
"f":0.
|
| 774 |
},
|
| 775 |
"conj":{
|
| 776 |
-
"p":0.
|
| 777 |
-
"r":0.
|
| 778 |
-
"f":0.
|
| 779 |
},
|
| 780 |
"cc":{
|
| 781 |
-
"p":0.
|
| 782 |
-
"r":0.
|
| 783 |
-
"f":0.
|
| 784 |
},
|
| 785 |
"obj":{
|
| 786 |
-
"p":0.
|
| 787 |
-
"r":0.
|
| 788 |
-
"f":0.
|
| 789 |
},
|
| 790 |
"obl:agent":{
|
| 791 |
"p":0.7941176471,
|
|
@@ -793,59 +793,64 @@
|
|
| 793 |
"f":0.84375
|
| 794 |
},
|
| 795 |
"acl:relcl":{
|
| 796 |
-
"p":0.
|
| 797 |
-
"r":0.
|
| 798 |
-
"f":0.
|
| 799 |
},
|
| 800 |
"mark":{
|
| 801 |
-
"p":0.
|
| 802 |
-
"r":0.
|
| 803 |
-
"f":0.
|
| 804 |
},
|
| 805 |
"advcl":{
|
| 806 |
-
"p":0.
|
| 807 |
-
"r":0.
|
| 808 |
-
"f":0.
|
| 809 |
},
|
| 810 |
"xcomp":{
|
| 811 |
-
"p":0.
|
| 812 |
-
"r":0.
|
| 813 |
-
"f":0.
|
| 814 |
},
|
| 815 |
"iobj":{
|
| 816 |
-
"p":0.
|
| 817 |
-
"r":0.
|
| 818 |
-
"f":0.
|
| 819 |
},
|
| 820 |
"appos":{
|
| 821 |
-
"p":0.
|
| 822 |
-
"r":0.
|
| 823 |
-
"f":0.
|
| 824 |
},
|
| 825 |
"fixed":{
|
| 826 |
-
"p":0.
|
| 827 |
-
"r":0.
|
| 828 |
-
"f":0.
|
| 829 |
},
|
| 830 |
"nummod":{
|
| 831 |
-
"p":0.
|
| 832 |
-
"r":0.
|
| 833 |
-
"f":0.
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 834 |
},
|
| 835 |
"aux":{
|
| 836 |
-
"p":0.
|
| 837 |
-
"r":0.
|
| 838 |
-
"f":0.
|
| 839 |
},
|
| 840 |
"csubj":{
|
| 841 |
-
"p":0.
|
| 842 |
-
"r":0.
|
| 843 |
-
"f":0.
|
| 844 |
},
|
| 845 |
"ccomp":{
|
| 846 |
-
"p":0.
|
| 847 |
-
"r":0.
|
| 848 |
-
"f":0.
|
| 849 |
},
|
| 850 |
"orphan":{
|
| 851 |
"p":0.0,
|
|
@@ -859,18 +864,18 @@
|
|
| 859 |
},
|
| 860 |
"aux:pass":{
|
| 861 |
"p":1.0,
|
| 862 |
-
"r":0.
|
| 863 |
-
"f":0.
|
| 864 |
},
|
| 865 |
"nsubj:pass":{
|
| 866 |
-
"p":0.
|
| 867 |
-
"r":0.
|
| 868 |
-
"f":0.
|
| 869 |
},
|
| 870 |
"parataxis":{
|
| 871 |
-
"p":0.
|
| 872 |
-
"r":0.
|
| 873 |
-
"f":0.
|
| 874 |
},
|
| 875 |
"list":{
|
| 876 |
"p":0.0,
|
|
@@ -878,19 +883,14 @@
|
|
| 878 |
"f":0.0
|
| 879 |
},
|
| 880 |
"expl":{
|
| 881 |
-
"p":0.
|
| 882 |
-
"r":
|
| 883 |
-
"f":0.
|
| 884 |
},
|
| 885 |
"compound":{
|
| 886 |
-
"p":0.
|
| 887 |
-
"r":0.
|
| 888 |
-
"f":0.
|
| 889 |
-
},
|
| 890 |
-
"dep":{
|
| 891 |
-
"p":0.0,
|
| 892 |
-
"r":0.0,
|
| 893 |
-
"f":0.0
|
| 894 |
},
|
| 895 |
"vocative":{
|
| 896 |
"p":0.0,
|
|
@@ -898,9 +898,9 @@
|
|
| 898 |
"f":0.0
|
| 899 |
},
|
| 900 |
"discourse":{
|
| 901 |
-
"p":
|
| 902 |
-
"r":0.
|
| 903 |
-
"f":0.
|
| 904 |
},
|
| 905 |
"expl:pass":{
|
| 906 |
"p":0.0,
|
|
@@ -913,34 +913,34 @@
|
|
| 913 |
"f":0.8571428571
|
| 914 |
}
|
| 915 |
},
|
| 916 |
-
"lemma_acc":0.
|
| 917 |
-
"tag_acc":0.
|
| 918 |
-
"ents_p":0.
|
| 919 |
-
"ents_r":0.
|
| 920 |
-
"ents_f":0.
|
| 921 |
"ents_per_type":{
|
| 922 |
"LOC":{
|
| 923 |
-
"p":0.
|
| 924 |
-
"r":0.
|
| 925 |
-
"f":0.
|
| 926 |
},
|
| 927 |
"PER":{
|
| 928 |
-
"p":0.
|
| 929 |
-
"r":0.
|
| 930 |
-
"f":0.
|
| 931 |
},
|
| 932 |
"ORG":{
|
| 933 |
-
"p":0.
|
| 934 |
-
"r":0.
|
| 935 |
-
"f":0.
|
| 936 |
},
|
| 937 |
"MISC":{
|
| 938 |
-
"p":0.
|
| 939 |
-
"r":0.
|
| 940 |
-
"f":0.
|
| 941 |
}
|
| 942 |
},
|
| 943 |
-
"speed":
|
| 944 |
},
|
| 945 |
"sources":[
|
| 946 |
{
|
|
|
|
| 1 |
{
|
| 2 |
"lang":"pt",
|
| 3 |
"name":"core_news_md",
|
| 4 |
+
"version":"3.5.0",
|
| 5 |
"description":"Portuguese pipeline optimized for CPU. Components: tok2vec, morphologizer, parser, lemmatizer (trainable_lemmatizer), senter, ner, attribute_ruler.",
|
| 6 |
"author":"Explosion",
|
| 7 |
"email":"contact@explosion.ai",
|
| 8 |
"url":"https://explosion.ai",
|
| 9 |
"license":"CC BY-SA 4.0",
|
| 10 |
+
"spacy_version":">=3.5.0,<3.6.0",
|
| 11 |
+
"spacy_git_version":"9e0322de1",
|
| 12 |
"vectors":{
|
| 13 |
"width":300,
|
| 14 |
"vectors":20000,
|
|
|
|
| 644 |
"token_p":0.9988117635,
|
| 645 |
"token_r":0.9995045581,
|
| 646 |
"token_f":0.9991580407,
|
| 647 |
+
"pos_acc":0.9702313141,
|
| 648 |
+
"morph_acc":0.9534422982,
|
| 649 |
+
"morph_micro_p":0.978440699,
|
| 650 |
+
"morph_micro_r":0.9751133553,
|
| 651 |
+
"morph_micro_f":0.9767741935,
|
| 652 |
"morph_per_feat":{
|
| 653 |
"Mood":{
|
| 654 |
+
"p":0.9770808203,
|
| 655 |
"r":0.9794437727,
|
| 656 |
+
"f":0.9782608696
|
| 657 |
},
|
| 658 |
"Number":{
|
| 659 |
+
"p":0.9926424546,
|
| 660 |
+
"r":0.9875408815,
|
| 661 |
+
"f":0.9900850964
|
| 662 |
},
|
| 663 |
"Person":{
|
| 664 |
+
"p":0.9712765957,
|
| 665 |
+
"r":0.9743863394,
|
| 666 |
+
"f":0.9728289824
|
| 667 |
},
|
| 668 |
"Tense":{
|
| 669 |
+
"p":0.9304677623,
|
| 670 |
+
"r":0.9684210526,
|
| 671 |
+
"f":0.9490651193
|
| 672 |
},
|
| 673 |
"VerbForm":{
|
| 674 |
+
"p":0.9732072498,
|
| 675 |
+
"r":0.9786053883,
|
| 676 |
+
"f":0.9758988542
|
| 677 |
},
|
| 678 |
"Gender":{
|
| 679 |
+
"p":0.9659508609,
|
| 680 |
+
"r":0.9548670874,
|
| 681 |
+
"f":0.9603769956
|
| 682 |
},
|
| 683 |
"PronType":{
|
| 684 |
+
"p":0.9890625,
|
| 685 |
+
"r":0.9813953488,
|
| 686 |
+
"f":0.9852140078
|
| 687 |
},
|
| 688 |
"Definite":{
|
| 689 |
+
"p":0.9952574526,
|
| 690 |
+
"r":0.9952574526,
|
| 691 |
+
"f":0.9952574526
|
| 692 |
},
|
| 693 |
"NumType":{
|
| 694 |
+
"p":0.9671052632,
|
| 695 |
+
"r":0.9639344262,
|
| 696 |
+
"f":0.9655172414
|
| 697 |
},
|
| 698 |
"Voice":{
|
| 699 |
+
"p":0.9195402299,
|
| 700 |
"r":0.9302325581,
|
| 701 |
+
"f":0.9248554913
|
| 702 |
},
|
| 703 |
"Polarity":{
|
| 704 |
+
"p":1.0,
|
| 705 |
+
"r":0.9861111111,
|
| 706 |
+
"f":0.993006993
|
| 707 |
},
|
| 708 |
"Case":{
|
| 709 |
+
"p":0.8571428571,
|
| 710 |
+
"r":0.8571428571,
|
| 711 |
+
"f":0.8571428571
|
| 712 |
}
|
| 713 |
},
|
| 714 |
+
"sents_p":0.9346642468,
|
| 715 |
+
"sents_r":0.9597677922,
|
| 716 |
+
"sents_f":0.9457981824,
|
| 717 |
+
"dep_uas":0.9013711257,
|
| 718 |
+
"dep_las":0.8593155894,
|
| 719 |
"dep_las_per_type":{
|
| 720 |
"cop":{
|
| 721 |
+
"p":0.8684210526,
|
| 722 |
"r":0.9361702128,
|
| 723 |
+
"f":0.9010238908
|
| 724 |
},
|
| 725 |
"root":{
|
| 726 |
+
"p":0.9364791289,
|
| 727 |
+
"r":0.9214285714,
|
| 728 |
+
"f":0.9288928893
|
| 729 |
},
|
| 730 |
"det":{
|
| 731 |
+
"p":0.9638554217,
|
| 732 |
+
"r":0.9659714599,
|
| 733 |
+
"f":0.9649122807
|
| 734 |
},
|
| 735 |
"amod":{
|
| 736 |
+
"p":0.9423558897,
|
| 737 |
+
"r":0.9170731707,
|
| 738 |
+
"f":0.9295426452
|
| 739 |
},
|
| 740 |
"nsubj":{
|
| 741 |
+
"p":0.8890814558,
|
| 742 |
+
"r":0.8952879581,
|
| 743 |
+
"f":0.892173913
|
| 744 |
},
|
| 745 |
"case":{
|
| 746 |
+
"p":0.973265436,
|
| 747 |
+
"r":0.9757498405,
|
| 748 |
+
"f":0.9745060548
|
| 749 |
},
|
| 750 |
"nmod":{
|
| 751 |
+
"p":0.8051948052,
|
| 752 |
+
"r":0.8285077951,
|
| 753 |
+
"f":0.8166849616
|
| 754 |
},
|
| 755 |
"flat:name":{
|
| 756 |
+
"p":0.8732394366,
|
| 757 |
+
"r":0.9018181818,
|
| 758 |
+
"f":0.8872987478
|
| 759 |
},
|
| 760 |
"acl":{
|
| 761 |
+
"p":0.6476190476,
|
| 762 |
+
"r":0.5862068966,
|
| 763 |
+
"f":0.6153846154
|
| 764 |
},
|
| 765 |
"advmod":{
|
| 766 |
+
"p":0.8163841808,
|
| 767 |
+
"r":0.8210227273,
|
| 768 |
+
"f":0.8186968839
|
| 769 |
},
|
| 770 |
"obl":{
|
| 771 |
+
"p":0.6811881188,
|
| 772 |
+
"r":0.70781893,
|
| 773 |
+
"f":0.6942482341
|
| 774 |
},
|
| 775 |
"conj":{
|
| 776 |
+
"p":0.5871212121,
|
| 777 |
+
"r":0.5719557196,
|
| 778 |
+
"f":0.5794392523
|
| 779 |
},
|
| 780 |
"cc":{
|
| 781 |
+
"p":0.8808510638,
|
| 782 |
+
"r":0.9,
|
| 783 |
+
"f":0.8903225806
|
| 784 |
},
|
| 785 |
"obj":{
|
| 786 |
+
"p":0.8734693878,
|
| 787 |
+
"r":0.8246628131,
|
| 788 |
+
"f":0.8483647175
|
| 789 |
},
|
| 790 |
"obl:agent":{
|
| 791 |
"p":0.7941176471,
|
|
|
|
| 793 |
"f":0.84375
|
| 794 |
},
|
| 795 |
"acl:relcl":{
|
| 796 |
+
"p":0.610619469,
|
| 797 |
+
"r":0.6509433962,
|
| 798 |
+
"f":0.6301369863
|
| 799 |
},
|
| 800 |
"mark":{
|
| 801 |
+
"p":0.8883495146,
|
| 802 |
+
"r":0.8551401869,
|
| 803 |
+
"f":0.8714285714
|
| 804 |
},
|
| 805 |
"advcl":{
|
| 806 |
+
"p":0.572519084,
|
| 807 |
+
"r":0.6696428571,
|
| 808 |
+
"f":0.6172839506
|
| 809 |
},
|
| 810 |
"xcomp":{
|
| 811 |
+
"p":0.8482142857,
|
| 812 |
+
"r":0.76,
|
| 813 |
+
"f":0.8016877637
|
| 814 |
},
|
| 815 |
"iobj":{
|
| 816 |
+
"p":0.6666666667,
|
| 817 |
+
"r":0.1818181818,
|
| 818 |
+
"f":0.2857142857
|
| 819 |
},
|
| 820 |
"appos":{
|
| 821 |
+
"p":0.5895953757,
|
| 822 |
+
"r":0.6219512195,
|
| 823 |
+
"f":0.6053412463
|
| 824 |
},
|
| 825 |
"fixed":{
|
| 826 |
+
"p":0.7692307692,
|
| 827 |
+
"r":0.7142857143,
|
| 828 |
+
"f":0.7407407407
|
| 829 |
},
|
| 830 |
"nummod":{
|
| 831 |
+
"p":0.9671052632,
|
| 832 |
+
"r":0.9423076923,
|
| 833 |
+
"f":0.9545454545
|
| 834 |
+
},
|
| 835 |
+
"dep":{
|
| 836 |
+
"p":0.0,
|
| 837 |
+
"r":0.0,
|
| 838 |
+
"f":0.0
|
| 839 |
},
|
| 840 |
"aux":{
|
| 841 |
+
"p":0.9076923077,
|
| 842 |
+
"r":0.9516129032,
|
| 843 |
+
"f":0.9291338583
|
| 844 |
},
|
| 845 |
"csubj":{
|
| 846 |
+
"p":0.6666666667,
|
| 847 |
+
"r":0.5,
|
| 848 |
+
"f":0.5714285714
|
| 849 |
},
|
| 850 |
"ccomp":{
|
| 851 |
+
"p":0.75,
|
| 852 |
+
"r":0.65,
|
| 853 |
+
"f":0.6964285714
|
| 854 |
},
|
| 855 |
"orphan":{
|
| 856 |
"p":0.0,
|
|
|
|
| 864 |
},
|
| 865 |
"aux:pass":{
|
| 866 |
"p":1.0,
|
| 867 |
+
"r":0.9705882353,
|
| 868 |
+
"f":0.9850746269
|
| 869 |
},
|
| 870 |
"nsubj:pass":{
|
| 871 |
+
"p":0.8679245283,
|
| 872 |
+
"r":0.8518518519,
|
| 873 |
+
"f":0.8598130841
|
| 874 |
},
|
| 875 |
"parataxis":{
|
| 876 |
+
"p":0.3818181818,
|
| 877 |
+
"r":0.2957746479,
|
| 878 |
+
"f":0.3333333333
|
| 879 |
},
|
| 880 |
"list":{
|
| 881 |
"p":0.0,
|
|
|
|
| 883 |
"f":0.0
|
| 884 |
},
|
| 885 |
"expl":{
|
| 886 |
+
"p":0.76,
|
| 887 |
+
"r":1.0,
|
| 888 |
+
"f":0.8636363636
|
| 889 |
},
|
| 890 |
"compound":{
|
| 891 |
+
"p":0.75,
|
| 892 |
+
"r":0.6,
|
| 893 |
+
"f":0.6666666667
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
| 894 |
},
|
| 895 |
"vocative":{
|
| 896 |
"p":0.0,
|
|
|
|
| 898 |
"f":0.0
|
| 899 |
},
|
| 900 |
"discourse":{
|
| 901 |
+
"p":1.0,
|
| 902 |
+
"r":0.3333333333,
|
| 903 |
+
"f":0.5
|
| 904 |
},
|
| 905 |
"expl:pass":{
|
| 906 |
"p":0.0,
|
|
|
|
| 913 |
"f":0.8571428571
|
| 914 |
}
|
| 915 |
},
|
| 916 |
+
"lemma_acc":0.9703333168,
|
| 917 |
+
"tag_acc":0.8962306206,
|
| 918 |
+
"ents_p":0.8937657035,
|
| 919 |
+
"ents_r":0.8960699034,
|
| 920 |
+
"ents_f":0.8949163203,
|
| 921 |
"ents_per_type":{
|
| 922 |
"LOC":{
|
| 923 |
+
"p":0.9114664452,
|
| 924 |
+
"r":0.9259845288,
|
| 925 |
+
"f":0.9186681318
|
| 926 |
},
|
| 927 |
"PER":{
|
| 928 |
+
"p":0.90687251,
|
| 929 |
+
"r":0.9220253165,
|
| 930 |
+
"f":0.9143861411
|
| 931 |
},
|
| 932 |
"ORG":{
|
| 933 |
+
"p":0.8461361015,
|
| 934 |
+
"r":0.8192986375,
|
| 935 |
+
"f":0.8325011348
|
| 936 |
},
|
| 937 |
"MISC":{
|
| 938 |
+
"p":0.8121030346,
|
| 939 |
+
"r":0.7613298048,
|
| 940 |
+
"f":0.785897217
|
| 941 |
}
|
| 942 |
},
|
| 943 |
+
"speed":11483.8429749162
|
| 944 |
},
|
| 945 |
"sources":[
|
| 946 |
{
|
morphologizer/model
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 213842
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:3b9f5a5f481e6192371ed4db0b5abf8cfb435a30a20a124866c57ccf7e223f64
|
| 3 |
size 213842
|
ner/model
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
-
size
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:2de62c3083e8eba77efadfa0a422d56cb91c882b2891c753e576eb6bec550866
|
| 3 |
+
size 6366382
|
parser/model
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 312369
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:30e6a46316ed20205a65f28a3c9f6432591f5289961958a076871ef6fa30e318
|
| 3 |
size 312369
|
pt_core_news_md-any-py3-none-any.whl
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
-
size
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:660d4a2ee1ccc513f96711245e4bdf816a6ac85121296038587f67a9ad32f17e
|
| 3 |
+
size 42370575
|
senter/model
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
size 219953
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:8ec07a20c6a4429942aab7cc5327af0f5b5167cfb6476cc70ca1d6c364087f12
|
| 3 |
size 219953
|
tok2vec/model
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
-
size
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:e99573d4420e527dbea54a4049897d0e7cba5c6da60ad6fd7a9f02c9aa07939e
|
| 3 |
+
size 6495793
|
tokenizer
CHANGED
|
@@ -1,3 +1,4 @@
|
|
| 1 |
-
��prefix_search��^\w{1,3}\$|^§|^%|^=|^—|^–|^\+(?![0-9])|^…|^……|^,|^:|^;|^\!|^\?|^¿|^؟|^¡|^\(|^\)|^\[|^\]|^\{|^\}|^<|^>|^_|^#|^\*|^&|^。|^?|^!|^,|^、|^;|^:|^~|^·|^।|^،|^۔|^؛|^٪|^\.\.+|^…|^\'|^"|^”|^“|^`|^‘|^´|^’|^‚|^,|^„|^»|^«|^「|^」|^『|^』|^(|^)|^〔|^〕|^【|^】|^《|^》|^〈|^〉|^\$|^£|^€|^¥|^฿|^US\$|^C\$|^A\$|^₽|^﷼|^₴|^₠|^₡|^₢|^₣|^₤|^₥|^₦|^₧|^₨|^₩|^₪|^₫|^€|^₭|^₮|^₯|^₰|^₱|^₲|^₳|^₴|^₵|^₶|^₷|^₸|^₹|^₺|^₻|^₼|^₽|^₾|^₿|^[\u00A6\u00A9\u00AE\u00B0\u0482\u058D\u058E\u060E\u060F\u06DE\u06E9\u06FD\u06FE\u07F6\u09FA\u0B70\u0BF3-\u0BF8\u0BFA\u0C7F\u0D4F\u0D79\u0F01-\u0F03\u0F13\u0F15-\u0F17\u0F1A-\u0F1F\u0F34\u0F36\u0F38\u0FBE-\u0FC5\u0FC7-\u0FCC\u0FCE\u0FCF\u0FD5-\u0FD8\u109E\u109F\u1390-\u1399\u1940\u19DE-\u19FF\u1B61-\u1B6A\u1B74-\u1B7C\u2100\u2101\u2103-\u2106\u2108\u2109\u2114\u2116\u2117\u211E-\u2123\u2125\u2127\u2129\u212E\u213A\u213B\u214A\u214C\u214D\u214F\u218A\u218B\u2195-\u2199\u219C-\u219F\u21A1\u21A2\u21A4\u21A5\u21A7-\u21AD\u21AF-\u21CD\u21D0\u21D1\u21D3\u21D5-\u21F3\u2300-\u2307\u230C-\u231F\u2322-\u2328\u232B-\u237B\u237D-\u239A\u23B4-\u23DB\u23E2-\u2426\u2440-\u244A\u249C-\u24E9\u2500-\u25B6\u25B8-\u25C0\u25C2-\u25F7\u2600-\u266E\u2670-\u2767\u2794-\u27BF\u2800-\u28FF\u2B00-\u2B2F\u2B45\u2B46\u2B4D-\u2B73\u2B76-\u2B95\u2B98-\u2BC8\u2BCA-\u2BFE\u2CE5-\u2CEA\u2E80-\u2E99\u2E9B-\u2EF3\u2F00-\u2FD5\u2FF0-\u2FFB\u3004\u3012\u3013\u3020\u3036\u3037\u303E\u303F\u3190\u3191\u3196-\u319F\u31C0-\u31E3\u3200-\u321E\u322A-\u3247\u3250\u3260-\u327F\u328A-\u32B0\u32C0-\u32FE\u3300-\u33FF\u4DC0-\u4DFF\uA490-\uA4C6\uA828-\uA82B\uA836\uA837\uA839\uAA77-\uAA79\uFDFD\uFFE4\uFFE8\uFFED\uFFEE\uFFFC\uFFFD\U00010137-\U0001013F\U00010179-\U00010189\U0001018C-\U0001018E\U00010190-\U0001019B\U000101A0\U000101D0-\U000101FC\U00010877\U00010878\U00010AC8\U0001173F\U00016B3C-\U00016B3F\U00016B45\U0001BC9C\U0001D000-\U0001D0F5\U0001D100-\U0001D126\U0001D129-\U0001D164\U0001D16A-\U0001D16C\U0001D183\U0001D184\U0001D18C-\U0001D1A9\U0001D1AE-\U0001D1E8\U0001D200-\U0001D241\U0001D245\U0001D300-\U0001D356\U0001D800-\U0001D9FF\U0001DA37-\U0001DA3A\U0001DA6D-\U0001DA74\U0001DA76-\U0001DA83\U0001DA85\U0001DA86\U0001ECAC\U0001F000-\U0001F02B\U0001F030-\U0001F093\U0001F0A0-\U0001F0AE\U0001F0B1-\U0001F0BF\U0001F0C1-\U0001F0CF\U0001F0D1-\U0001F0F5\U0001F110-\U0001F16B\U0001F170-\U0001F1AC\U0001F1E6-\U0001F202\U0001F210-\U0001F23B\U0001F240-\U0001F248\U0001F250\U0001F251\U0001F260-\U0001F265\U0001F300-\U0001F3FA\U0001F400-\U0001F6D4\U0001F6E0-\U0001F6EC\U0001F6F0-\U0001F6F9\U0001F700-\U0001F773\U0001F780-\U0001F7D8\U0001F800-\U0001F80B\U0001F810-\U0001F847\U0001F850-\U0001F859\U0001F860-\U0001F887\U0001F890-\U0001F8AD\U0001F900-\U0001F90B\U0001F910-\U0001F93E\U0001F940-\U0001F970\U0001F973-\U0001F976\U0001F97A\U0001F97C-\U0001F9A2\U0001F9B0-\U0001F9B9\U0001F9C0-\U0001F9C2\U0001F9D0-\U0001F9FF\U0001FA60-\U0001FA6D]�suffix_search�2y…$|……$|,$|:$|;$|\!$|\?$|¿$|؟$|¡$|\($|\)$|\[$|\]$|\{$|\}$|<$|>$|_$|#$|\*$|&$|。$|?$|!$|,$|、$|;$|:$|~$|·$|।$|،$|۔$|؛$|٪$|\.\.+$|…$|\'$|"$|”$|“$|`$|‘$|´$|’$|‚$|,$|„$|»$|«$|「$|」$|『$|』$|($|)$|〔$|〕$|【$|】$|《$|》$|〈$|〉$|[\u00A6\u00A9\u00AE\u00B0\u0482\u058D\u058E\u060E\u060F\u06DE\u06E9\u06FD\u06FE\u07F6\u09FA\u0B70\u0BF3-\u0BF8\u0BFA\u0C7F\u0D4F\u0D79\u0F01-\u0F03\u0F13\u0F15-\u0F17\u0F1A-\u0F1F\u0F34\u0F36\u0F38\u0FBE-\u0FC5\u0FC7-\u0FCC\u0FCE\u0FCF\u0FD5-\u0FD8\u109E\u109F\u1390-\u1399\u1940\u19DE-\u19FF\u1B61-\u1B6A\u1B74-\u1B7C\u2100\u2101\u2103-\u2106\u2108\u2109\u2114\u2116\u2117\u211E-\u2123\u2125\u2127\u2129\u212E\u213A\u213B\u214A\u214C\u214D\u214F\u218A\u218B\u2195-\u2199\u219C-\u219F\u21A1\u21A2\u21A4\u21A5\u21A7-\u21AD\u21AF-\u21CD\u21D0\u21D1\u21D3\u21D5-\u21F3\u2300-\u2307\u230C-\u231F\u2322-\u2328\u232B-\u237B\u237D-\u239A\u23B4-\u23DB\u23E2-\u2426\u2440-\u244A\u249C-\u24E9\u2500-\u25B6\u25B8-\u25C0\u25C2-\u25F7\u2600-\u266E\u2670-\u2767\u2794-\u27BF\u2800-\u28FF\u2B00-\u2B2F\u2B45\u2B46\u2B4D-\u2B73\u2B76-\u2B95\u2B98-\u2BC8\u2BCA-\u2BFE\u2CE5-\u2CEA\u2E80-\u2E99\u2E9B-\u2EF3\u2F00-\u2FD5\u2FF0-\u2FFB\u3004\u3012\u3013\u3020\u3036\u3037\u303E\u303F\u3190\u3191\u3196-\u319F\u31C0-\u31E3\u3200-\u321E\u322A-\u3247\u3250\u3260-\u327F\u328A-\u32B0\u32C0-\u32FE\u3300-\u33FF\u4DC0-\u4DFF\uA490-\uA4C6\uA828-\uA82B\uA836\uA837\uA839\uAA77-\uAA79\uFDFD\uFFE4\uFFE8\uFFED\uFFEE\uFFFC\uFFFD\U00010137-\U0001013F\U00010179-\U00010189\U0001018C-\U0001018E\U00010190-\U0001019B\U000101A0\U000101D0-\U000101FC\U00010877\U00010878\U00010AC8\U0001173F\U00016B3C-\U00016B3F\U00016B45\U0001BC9C\U0001D000-\U0001D0F5\U0001D100-\U0001D126\U0001D129-\U0001D164\U0001D16A-\U0001D16C\U0001D183\U0001D184\U0001D18C-\U0001D1A9\U0001D1AE-\U0001D1E8\U0001D200-\U0001D241\U0001D245\U0001D300-\U0001D356\U0001D800-\U0001D9FF\U0001DA37-\U0001DA3A\U0001DA6D-\U0001DA74\U0001DA76-\U0001DA83\U0001DA85\U0001DA86\U0001ECAC\U0001F000-\U0001F02B\U0001F030-\U0001F093\U0001F0A0-\U0001F0AE\U0001F0B1-\U0001F0BF\U0001F0C1-\U0001F0CF\U0001F0D1-\U0001F0F5\U0001F110-\U0001F16B\U0001F170-\U0001F1AC\U0001F1E6-\U0001F202\U0001F210-\U0001F23B\U0001F240-\U0001F248\U0001F250\U0001F251\U0001F260-\U0001F265\U0001F300-\U0001F3FA\U0001F400-\U0001F6D4\U0001F6E0-\U0001F6EC\U0001F6F0-\U0001F6F9\U0001F700-\U0001F773\U0001F780-\U0001F7D8\U0001F800-\U0001F80B\U0001F810-\U0001F847\U0001F850-\U0001F859\U0001F860-\U0001F887\U0001F890-\U0001F8AD\U0001F900-\U0001F90B\U0001F910-\U0001F93E\U0001F940-\U0001F970\U0001F973-\U0001F976\U0001F97A\U0001F97C-\U0001F9A2\U0001F9B0-\U0001F9B9\U0001F9C0-\U0001F9C2\U0001F9D0-\U0001F9FF\U0001FA60-\U0001FA6D]$|'s$|'S$|’s$|’S$|—$|–$|(?<=[0-9])\+$|(?<=°[FfCcKk])\.$|(?<=[0-9])(?:\$|£|€|¥|฿|US\$|C\$|A\$|₽|﷼|₴|₠|₡|₢|₣|₤|₥|₦|₧|₨|₩|₪|₫|€|₭|₮|₯|₰|₱|₲|₳|₴|₵|₶|₷|₸|₹|₺|₻|₼|₽|₾|₿)$|(?<=[0-9])(?:km|km²|km³|m|m²|m³|dm|dm²|dm³|cm|cm²|cm³|mm|mm²|mm³|ha|µm|nm|yd|in|ft|kg|g|mg|µg|t|lb|oz|m/s|km/h|kmh|mph|hPa|Pa|mbar|mb|MB|kb|KB|gb|GB|tb|TB|T|G|M|K|%|км|км²|км³|м|м²|м³|дм|дм²|дм³|см|см²|см³|мм|мм²|мм³|нм|кг|г|мг|м/с|км/ч|кПа|Па|мбар|Кб|КБ|кб|Мб|МБ|мб|Гб|ГБ|гб|Тб|ТБ|тбكم|كم²|كم³|م|م²|م³|سم|سم²|سم³|مم|مم²|مم³|كم|غرام|جرام|جم|كغ|ملغ|كوب|اكواب)$|(?<=[0-9a-z\uFF41-\uFF5A\u00DF-\u00F6\u00F8-\u00FF\u0101\u0103\u0105\u0107\u0109\u010B\u010D\u010F\u0111\u0113\u0115\u0117\u0119\u011B\u011D\u011F\u0121\u0123\u0125\u0127\u0129\u012B\u012D\u012F\u0131\u0133\u0135\u0137\u0138\u013A\u013C\u013E\u0140\u0142\u0144\u0146\u0148\u0149\u014B\u014D\u014F\u0151\u0153\u0155\u0157\u0159\u015B\u015D\u015F\u0161\u0163\u0165\u0167\u0169\u016B\u016D\u016F\u0171\u0173\u0175\u0177\u017A\u017C\u017E\u017F\u0180\u0183\u0185\u0188\u018C\u018D\u0192\u0195\u0199-\u019B\u019E\u01A1\u01A3\u01A5\u01A8\u01AA\u01AB\u01AD\u01B0\u01B4\u01B6\u01B9\u01BA\u01BD-\u01BF\u01C6\u01C9\u01CC\u01CE\u01D0\u01D2\u01D4\u01D6\u01D8\u01DA\u01DC\u01DD\u01DF\u01E1\u01E3\u01E5\u01E7\u01E9\u01EB\u01ED\u01EF\u01F0\u01F3\u01F5\u01F9\u01FB\u01FD\u01FF\u0201\u0203\u0205\u0207\u0209\u020B\u020D\u020F\u0211\u0213\u0215\u0217\u0219\u021B\u021D\u021F\u0221\u0223\u0225\u0227\u0229\u022B\u022D\u022F\u0231\u0233-\u0239\u023C\u023F\u0240\u0242\u0247\u0249\u024B\u024D\u024F\u2C61\u2C65\u2C66\u2C68\u2C6A\u2C6C\u2C71\u2C73\u2C74\u2C76-\u2C7B\uA723\uA725\uA727\uA729\uA72B\uA72D\uA72F-\uA731\uA733\uA735\uA737\uA739\uA73B\uA73D\uA73F\uA741\uA743\uA745\uA747\uA749\uA74B\uA74D\uA74F\uA751\uA753\uA755\uA757\uA759\uA75B\uA75D\uA75F\uA761\uA763\uA765\uA767\uA769\uA76B\uA76D\uA76F\uA771-\uA778\uA77A\uA77C\uA77F\uA781\uA783\uA785\uA787\uA78C\uA78E\uA791\uA793-\uA795\uA797\uA799\uA79B\uA79D\uA79F\uA7A1\uA7A3\uA7A5\uA7A7\uA7A9\uA7AF\uA7B5\uA7B7\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E01\u1E03\u1E05\u1E07\u1E09\u1E0B\u1E0D\u1E0F\u1E11\u1E13\u1E15\u1E17\u1E19\u1E1B\u1E1D\u1E1F\u1E21\u1E23\u1E25\u1E27\u1E29\u1E2B\u1E2D\u1E2F\u1E31\u1E33\u1E35\u1E37\u1E39\u1E3B\u1E3D\u1E3F\u1E41\u1E43\u1E45\u1E47\u1E49\u1E4B\u1E4D\u1E4F\u1E51\u1E53\u1E55\u1E57\u1E59\u1E5B\u1E5D\u1E5F\u1E61\u1E63\u1E65\u1E67\u1E69\u1E6B\u1E6D\u1E6F\u1E71\u1E73\u1E75\u1E77\u1E79\u1E7B\u1E7D\u1E7F\u1E81\u1E83\u1E85\u1E87\u1E89\u1E8B\u1E8D\u1E8F\u1E91\u1E93\u1E95-\u1E9D\u1E9F\u1EA1\u1EA3\u1EA5\u1EA7\u1EA9\u1EAB\u1EAD\u1EAF\u1EB1\u1EB3\u1EB5\u1EB7\u1EB9\u1EBB\u1EBD\u1EBF\u1EC1\u1EC3\u1EC5\u1EC7\u1EC9\u1ECB\u1ECD\u1ECF\u1ED1\u1ED3\u1ED5\u1ED7\u1ED9\u1EDB\u1EDD\u1EDF\u1EE1\u1EE3\u1EE5\u1EE7\u1EE9\u1EEB\u1EED\u1EEF\u1EF1\u1EF3\u1EF5\u1EF7\u1EF9\u1EFB\u1EFD\u1EFFёа-яәөүҗңһα-ωάέίόώήύа-щюяіїєґѓѕјљњќѐѝ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F%²\-\+…|……|,|:|;|\!|\?|¿|؟|¡|\(|\)|\[|\]|\{|\}|<|>|_|#|\*|&|。|?|!|,|、|;|:|~|·|।|،|۔|؛|٪(?:\'"”“`‘´’‚,„»«「」『』()〔〕【】《》〈〉)])\.$|(?<=[A-Z\uFF21-\uFF3A\u00C0-\u00D6\u00D8-\u00DE\u0100\u0102\u0104\u0106\u0108\u010A\u010C\u010E\u0110\u0112\u0114\u0116\u0118\u011A\u011C\u011E\u0120\u0122\u0124\u0126\u0128\u012A\u012C\u012E\u0130\u0132\u0134\u0136\u0139\u013B\u013D\u013F\u0141\u0143\u0145\u0147\u014A\u014C\u014E\u0150\u0152\u0154\u0156\u0158\u015A\u015C\u015E\u0160\u0162\u0164\u0166\u0168\u016A\u016C\u016E\u0170\u0172\u0174\u0176\u0178\u0179\u017B\u017D\u0181\u0182\u0184\u0186\u0187\u0189-\u018B\u018E-\u0191\u0193\u0194\u0196-\u0198\u019C\u019D\u019F\u01A0\u01A2\u01A4\u01A6\u01A7\u01A9\u01AC\u01AE\u01AF\u01B1-\u01B3\u01B5\u01B7\u01B8\u01BC\u01C4\u01C7\u01CA\u01CD\u01CF\u01D1\u01D3\u01D5\u01D7\u01D9\u01DB\u01DE\u01E0\u01E2\u01E4\u01E6\u01E8\u01EA\u01EC\u01EE\u01F1\u01F4\u01F6-\u01F8\u01FA\u01FC\u01FE\u0200\u0202\u0204\u0206\u0208\u020A\u020C\u020E\u0210\u0212\u0214\u0216\u0218\u021A\u021C\u021E\u0220\u0222\u0224\u0226\u0228\u022A\u022C\u022E\u0230\u0232\u023A\u023B\u023D\u023E\u0241\u0243-\u0246\u0248\u024A\u024C\u024E\u2C60\u2C62-\u2C64\u2C67\u2C69\u2C6B\u2C6D-\u2C70\u2C72\u2C75\u2C7E\u2C7F\uA722\uA724\uA726\uA728\uA72A\uA72C\uA72E\uA732\uA734\uA736\uA738\uA73A\uA73C\uA73E\uA740\uA742\uA744\uA746\uA748\uA74A\uA74C\uA74E\uA750\uA752\uA754\uA756\uA758\uA75A\uA75C\uA75E\uA760\uA762\uA764\uA766\uA768\uA76A\uA76C\uA76E\uA779\uA77B\uA77D\uA77E\uA780\uA782\uA784\uA786\uA78B\uA78D\uA790\uA792\uA796\uA798\uA79A\uA79C\uA79E\uA7A0\uA7A2\uA7A4\uA7A6\uA7A8\uA7AA-\uA7AE\uA7B0-\uA7B4\uA7B6\uA7B8\u1E00\u1E02\u1E04\u1E06\u1E08\u1E0A\u1E0C\u1E0E\u1E10\u1E12\u1E14\u1E16\u1E18\u1E1A\u1E1C\u1E1E\u1E20\u1E22\u1E24\u1E26\u1E28\u1E2A\u1E2C\u1E2E\u1E30\u1E32\u1E34\u1E36\u1E38\u1E3A\u1E3C\u1E3E\u1E40\u1E42\u1E44\u1E46\u1E48\u1E4A\u1E4C\u1E4E\u1E50\u1E52\u1E54\u1E56\u1E58\u1E5A\u1E5C\u1E5E\u1E60\u1E62\u1E64\u1E66\u1E68\u1E6A\u1E6C\u1E6E\u1E70\u1E72\u1E74\u1E76\u1E78\u1E7A\u1E7C\u1E7E\u1E80\u1E82\u1E84\u1E86\u1E88\u1E8A\u1E8C\u1E8E\u1E90\u1E92\u1E94\u1E9E\u1EA0\u1EA2\u1EA4\u1EA6\u1EA8\u1EAA\u1EAC\u1EAE\u1EB0\u1EB2\u1EB4\u1EB6\u1EB8\u1EBA\u1EBC\u1EBE\u1EC0\u1EC2\u1EC4\u1EC6\u1EC8\u1ECA\u1ECC\u1ECE\u1ED0\u1ED2\u1ED4\u1ED6\u1ED8\u1EDA\u1EDC\u1EDE\u1EE0\u1EE2\u1EE4\u1EE6\u1EE8\u1EEA\u1EEC\u1EEE\u1EF0\u1EF2\u1EF4\u1EF6\u1EF8\u1EFA\u1EFC\u1EFEЁА-ЯӘӨҮҖҢҺΑ-ΩΆΈΊΌΏΉΎА-ЩЮЯІЇЄҐЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F][A-Z\uFF21-\uFF3A\u00C0-\u00D6\u00D8-\u00DE\u0100\u0102\u0104\u0106\u0108\u010A\u010C\u010E\u0110\u0112\u0114\u0116\u0118\u011A\u011C\u011E\u0120\u0122\u0124\u0126\u0128\u012A\u012C\u012E\u0130\u0132\u0134\u0136\u0139\u013B\u013D\u013F\u0141\u0143\u0145\u0147\u014A\u014C\u014E\u0150\u0152\u0154\u0156\u0158\u015A\u015C\u015E\u0160\u0162\u0164\u0166\u0168\u016A\u016C\u016E\u0170\u0172\u0174\u0176\u0178\u0179\u017B\u017D\u0181\u0182\u0184\u0186\u0187\u0189-\u018B\u018E-\u0191\u0193\u0194\u0196-\u0198\u019C\u019D\u019F\u01A0\u01A2\u01A4\u01A6\u01A7\u01A9\u01AC\u01AE\u01AF\u01B1-\u01B3\u01B5\u01B7\u01B8\u01BC\u01C4\u01C7\u01CA\u01CD\u01CF\u01D1\u01D3\u01D5\u01D7\u01D9\u01DB\u01DE\u01E0\u01E2\u01E4\u01E6\u01E8\u01EA\u01EC\u01EE\u01F1\u01F4\u01F6-\u01F8\u01FA\u01FC\u01FE\u0200\u0202\u0204\u0206\u0208\u020A\u020C\u020E\u0210\u0212\u0214\u0216\u0218\u021A\u021C\u021E\u0220\u0222\u0224\u0226\u0228\u022A\u022C\u022E\u0230\u0232\u023A\u023B\u023D\u023E\u0241\u0243-\u0246\u0248\u024A\u024C\u024E\u2C60\u2C62-\u2C64\u2C67\u2C69\u2C6B\u2C6D-\u2C70\u2C72\u2C75\u2C7E\u2C7F\uA722\uA724\uA726\uA728\uA72A\uA72C\uA72E\uA732\uA734\uA736\uA738\uA73A\uA73C\uA73E\uA740\uA742\uA744\uA746\uA748\uA74A\uA74C\uA74E\uA750\uA752\uA754\uA756\uA758\uA75A\uA75C\uA75E\uA760\uA762\uA764\uA766\uA768\uA76A\uA76C\uA76E\uA779\uA77B\uA77D\uA77E\uA780\uA782\uA784\uA786\uA78B\uA78D\uA790\uA792\uA796\uA798\uA79A\uA79C\uA79E\uA7A0\uA7A2\uA7A4\uA7A6\uA7A8\uA7AA-\uA7AE\uA7B0-\uA7B4\uA7B6\uA7B8\u1E00\u1E02\u1E04\u1E06\u1E08\u1E0A\u1E0C\u1E0E\u1E10\u1E12\u1E14\u1E16\u1E18\u1E1A\u1E1C\u1E1E\u1E20\u1E22\u1E24\u1E26\u1E28\u1E2A\u1E2C\u1E2E\u1E30\u1E32\u1E34\u1E36\u1E38\u1E3A\u1E3C\u1E3E\u1E40\u1E42\u1E44\u1E46\u1E48\u1E4A\u1E4C\u1E4E\u1E50\u1E52\u1E54\u1E56\u1E58\u1E5A\u1E5C\u1E5E\u1E60\u1E62\u1E64\u1E66\u1E68\u1E6A\u1E6C\u1E6E\u1E70\u1E72\u1E74\u1E76\u1E78\u1E7A\u1E7C\u1E7E\u1E80\u1E82\u1E84\u1E86\u1E88\u1E8A\u1E8C\u1E8E\u1E90\u1E92\u1E94\u1E9E\u1EA0\u1EA2\u1EA4\u1EA6\u1EA8\u1EAA\u1EAC\u1EAE\u1EB0\u1EB2\u1EB4\u1EB6\u1EB8\u1EBA\u1EBC\u1EBE\u1EC0\u1EC2\u1EC4\u1EC6\u1EC8\u1ECA\u1ECC\u1ECE\u1ED0\u1ED2\u1ED4\u1ED6\u1ED8\u1EDA\u1EDC\u1EDE\u1EE0\u1EE2\u1EE4\u1EE6\u1EE8\u1EEA\u1EEC\u1EEE\u1EF0\u1EF2\u1EF4\u1EF6\u1EF8\u1EFA\u1EFC\u1EFEЁА-ЯӘӨҮҖҢҺΑ-ΩΆΈΊΌΏΉΎА-ЩЮЯІЇЄҐЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])\.$�infix_finditer�>�(\w+-\w+(-\w+)*)|\.\.+|…|[\u00A6\u00A9\u00AE\u00B0\u0482\u058D\u058E\u060E\u060F\u06DE\u06E9\u06FD\u06FE\u07F6\u09FA\u0B70\u0BF3-\u0BF8\u0BFA\u0C7F\u0D4F\u0D79\u0F01-\u0F03\u0F13\u0F15-\u0F17\u0F1A-\u0F1F\u0F34\u0F36\u0F38\u0FBE-\u0FC5\u0FC7-\u0FCC\u0FCE\u0FCF\u0FD5-\u0FD8\u109E\u109F\u1390-\u1399\u1940\u19DE-\u19FF\u1B61-\u1B6A\u1B74-\u1B7C\u2100\u2101\u2103-\u2106\u2108\u2109\u2114\u2116\u2117\u211E-\u2123\u2125\u2127\u2129\u212E\u213A\u213B\u214A\u214C\u214D\u214F\u218A\u218B\u2195-\u2199\u219C-\u219F\u21A1\u21A2\u21A4\u21A5\u21A7-\u21AD\u21AF-\u21CD\u21D0\u21D1\u21D3\u21D5-\u21F3\u2300-\u2307\u230C-\u231F\u2322-\u2328\u232B-\u237B\u237D-\u239A\u23B4-\u23DB\u23E2-\u2426\u2440-\u244A\u249C-\u24E9\u2500-\u25B6\u25B8-\u25C0\u25C2-\u25F7\u2600-\u266E\u2670-\u2767\u2794-\u27BF\u2800-\u28FF\u2B00-\u2B2F\u2B45\u2B46\u2B4D-\u2B73\u2B76-\u2B95\u2B98-\u2BC8\u2BCA-\u2BFE\u2CE5-\u2CEA\u2E80-\u2E99\u2E9B-\u2EF3\u2F00-\u2FD5\u2FF0-\u2FFB\u3004\u3012\u3013\u3020\u3036\u3037\u303E\u303F\u3190\u3191\u3196-\u319F\u31C0-\u31E3\u3200-\u321E\u322A-\u3247\u3250\u3260-\u327F\u328A-\u32B0\u32C0-\u32FE\u3300-\u33FF\u4DC0-\u4DFF\uA490-\uA4C6\uA828-\uA82B\uA836\uA837\uA839\uAA77-\uAA79\uFDFD\uFFE4\uFFE8\uFFED\uFFEE\uFFFC\uFFFD\U00010137-\U0001013F\U00010179-\U00010189\U0001018C-\U0001018E\U00010190-\U0001019B\U000101A0\U000101D0-\U000101FC\U00010877\U00010878\U00010AC8\U0001173F\U00016B3C-\U00016B3F\U00016B45\U0001BC9C\U0001D000-\U0001D0F5\U0001D100-\U0001D126\U0001D129-\U0001D164\U0001D16A-\U0001D16C\U0001D183\U0001D184\U0001D18C-\U0001D1A9\U0001D1AE-\U0001D1E8\U0001D200-\U0001D241\U0001D245\U0001D300-\U0001D356\U0001D800-\U0001D9FF\U0001DA37-\U0001DA3A\U0001DA6D-\U0001DA74\U0001DA76-\U0001DA83\U0001DA85\U0001DA86\U0001ECAC\U0001F000-\U0001F02B\U0001F030-\U0001F093\U0001F0A0-\U0001F0AE\U0001F0B1-\U0001F0BF\U0001F0C1-\U0001F0CF\U0001F0D1-\U0001F0F5\U0001F110-\U0001F16B\U0001F170-\U0001F1AC\U0001F1E6-\U0001F202\U0001F210-\U0001F23B\U0001F240-\U0001F248\U0001F250\U0001F251\U0001F260-\U0001F265\U0001F300-\U0001F3FA\U0001F400-\U0001F6D4\U0001F6E0-\U0001F6EC\U0001F6F0-\U0001F6F9\U0001F700-\U0001F773\U0001F780-\U0001F7D8\U0001F800-\U0001F80B\U0001F810-\U0001F847\U0001F850-\U0001F859\U0001F860-\U0001F887\U0001F890-\U0001F8AD\U0001F900-\U0001F90B\U0001F910-\U0001F93E\U0001F940-\U0001F970\U0001F973-\U0001F976\U0001F97A\U0001F97C-\U0001F9A2\U0001F9B0-\U0001F9B9\U0001F9C0-\U0001F9C2\U0001F9D0-\U0001F9FF\U0001FA60-\U0001FA6D]|(?<=[0-9])[+\-\*^](?=[0-9-])|(?<=[a-z\uFF41-\uFF5A\u00DF-\u00F6\u00F8-\u00FF\u0101\u0103\u0105\u0107\u0109\u010B\u010D\u010F\u0111\u0113\u0115\u0117\u0119\u011B\u011D\u011F\u0121\u0123\u0125\u0127\u0129\u012B\u012D\u012F\u0131\u0133\u0135\u0137\u0138\u013A\u013C\u013E\u0140\u0142\u0144\u0146\u0148\u0149\u014B\u014D\u014F\u0151\u0153\u0155\u0157\u0159\u015B\u015D\u015F\u0161\u0163\u0165\u0167\u0169\u016B\u016D\u016F\u0171\u0173\u0175\u0177\u017A\u017C\u017E\u017F\u0180\u0183\u0185\u0188\u018C\u018D\u0192\u0195\u0199-\u019B\u019E\u01A1\u01A3\u01A5\u01A8\u01AA\u01AB\u01AD\u01B0\u01B4\u01B6\u01B9\u01BA\u01BD-\u01BF\u01C6\u01C9\u01CC\u01CE\u01D0\u01D2\u01D4\u01D6\u01D8\u01DA\u01DC\u01DD\u01DF\u01E1\u01E3\u01E5\u01E7\u01E9\u01EB\u01ED\u01EF\u01F0\u01F3\u01F5\u01F9\u01FB\u01FD\u01FF\u0201\u0203\u0205\u0207\u0209\u020B\u020D\u020F\u0211\u0213\u0215\u0217\u0219\u021B\u021D\u021F\u0221\u0223\u0225\u0227\u0229\u022B\u022D\u022F\u0231\u0233-\u0239\u023C\u023F\u0240\u0242\u0247\u0249\u024B\u024D\u024F\u2C61\u2C65\u2C66\u2C68\u2C6A\u2C6C\u2C71\u2C73\u2C74\u2C76-\u2C7B\uA723\uA725\uA727\uA729\uA72B\uA72D\uA72F-\uA731\uA733\uA735\uA737\uA739\uA73B\uA73D\uA73F\uA741\uA743\uA745\uA747\uA749\uA74B\uA74D\uA74F\uA751\uA753\uA755\uA757\uA759\uA75B\uA75D\uA75F\uA761\uA763\uA765\uA767\uA769\uA76B\uA76D\uA76F\uA771-\uA778\uA77A\uA77C\uA77F\uA781\uA783\uA785\uA787\uA78C\uA78E\uA791\uA793-\uA795\uA797\uA799\uA79B\uA79D\uA79F\uA7A1\uA7A3\uA7A5\uA7A7\uA7A9\uA7AF\uA7B5\uA7B7\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E01\u1E03\u1E05\u1E07\u1E09\u1E0B\u1E0D\u1E0F\u1E11\u1E13\u1E15\u1E17\u1E19\u1E1B\u1E1D\u1E1F\u1E21\u1E23\u1E25\u1E27\u1E29\u1E2B\u1E2D\u1E2F\u1E31\u1E33\u1E35\u1E37\u1E39\u1E3B\u1E3D\u1E3F\u1E41\u1E43\u1E45\u1E47\u1E49\u1E4B\u1E4D\u1E4F\u1E51\u1E53\u1E55\u1E57\u1E59\u1E5B\u1E5D\u1E5F\u1E61\u1E63\u1E65\u1E67\u1E69\u1E6B\u1E6D\u1E6F\u1E71\u1E73\u1E75\u1E77\u1E79\u1E7B\u1E7D\u1E7F\u1E81\u1E83\u1E85\u1E87\u1E89\u1E8B\u1E8D\u1E8F\u1E91\u1E93\u1E95-\u1E9D\u1E9F\u1EA1\u1EA3\u1EA5\u1EA7\u1EA9\u1EAB\u1EAD\u1EAF\u1EB1\u1EB3\u1EB5\u1EB7\u1EB9\u1EBB\u1EBD\u1EBF\u1EC1\u1EC3\u1EC5\u1EC7\u1EC9\u1ECB\u1ECD\u1ECF\u1ED1\u1ED3\u1ED5\u1ED7\u1ED9\u1EDB\u1EDD\u1EDF\u1EE1\u1EE3\u1EE5\u1EE7\u1EE9\u1EEB\u1EED\u1EEF\u1EF1\u1EF3\u1EF5\u1EF7\u1EF9\u1EFB\u1EFD\u1EFFёа-яәөүҗңһα-ωάέίόώήύа-щюяіїєґѓѕјљњќѐѝ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F\'"”“`‘´’‚,„»«「」『』()〔〕【】《》〈〉])\.(?=[A-Z\uFF21-\uFF3A\u00C0-\u00D6\u00D8-\u00DE\u0100\u0102\u0104\u0106\u0108\u010A\u010C\u010E\u0110\u0112\u0114\u0116\u0118\u011A\u011C\u011E\u0120\u0122\u0124\u0126\u0128\u012A\u012C\u012E\u0130\u0132\u0134\u0136\u0139\u013B\u013D\u013F\u0141\u0143\u0145\u0147\u014A\u014C\u014E\u0150\u0152\u0154\u0156\u0158\u015A\u015C\u015E\u0160\u0162\u0164\u0166\u0168\u016A\u016C\u016E\u0170\u0172\u0174\u0176\u0178\u0179\u017B\u017D\u0181\u0182\u0184\u0186\u0187\u0189-\u018B\u018E-\u0191\u0193\u0194\u0196-\u0198\u019C\u019D\u019F\u01A0\u01A2\u01A4\u01A6\u01A7\u01A9\u01AC\u01AE\u01AF\u01B1-\u01B3\u01B5\u01B7\u01B8\u01BC\u01C4\u01C7\u01CA\u01CD\u01CF\u01D1\u01D3\u01D5\u01D7\u01D9\u01DB\u01DE\u01E0\u01E2\u01E4\u01E6\u01E8\u01EA\u01EC\u01EE\u01F1\u01F4\u01F6-\u01F8\u01FA\u01FC\u01FE\u0200\u0202\u0204\u0206\u0208\u020A\u020C\u020E\u0210\u0212\u0214\u0216\u0218\u021A\u021C\u021E\u0220\u0222\u0224\u0226\u0228\u022A\u022C\u022E\u0230\u0232\u023A\u023B\u023D\u023E\u0241\u0243-\u0246\u0248\u024A\u024C\u024E\u2C60\u2C62-\u2C64\u2C67\u2C69\u2C6B\u2C6D-\u2C70\u2C72\u2C75\u2C7E\u2C7F\uA722\uA724\uA726\uA728\uA72A\uA72C\uA72E\uA732\uA734\uA736\uA738\uA73A\uA73C\uA73E\uA740\uA742\uA744\uA746\uA748\uA74A\uA74C\uA74E\uA750\uA752\uA754\uA756\uA758\uA75A\uA75C\uA75E\uA760\uA762\uA764\uA766\uA768\uA76A\uA76C\uA76E\uA779\uA77B\uA77D\uA77E\uA780\uA782\uA784\uA786\uA78B\uA78D\uA790\uA792\uA796\uA798\uA79A\uA79C\uA79E\uA7A0\uA7A2\uA7A4\uA7A6\uA7A8\uA7AA-\uA7AE\uA7B0-\uA7B4\uA7B6\uA7B8\u1E00\u1E02\u1E04\u1E06\u1E08\u1E0A\u1E0C\u1E0E\u1E10\u1E12\u1E14\u1E16\u1E18\u1E1A\u1E1C\u1E1E\u1E20\u1E22\u1E24\u1E26\u1E28\u1E2A\u1E2C\u1E2E\u1E30\u1E32\u1E34\u1E36\u1E38\u1E3A\u1E3C\u1E3E\u1E40\u1E42\u1E44\u1E46\u1E48\u1E4A\u1E4C\u1E4E\u1E50\u1E52\u1E54\u1E56\u1E58\u1E5A\u1E5C\u1E5E\u1E60\u1E62\u1E64\u1E66\u1E68\u1E6A\u1E6C\u1E6E\u1E70\u1E72\u1E74\u1E76\u1E78\u1E7A\u1E7C\u1E7E\u1E80\u1E82\u1E84\u1E86\u1E88\u1E8A\u1E8C\u1E8E\u1E90\u1E92\u1E94\u1E9E\u1EA0\u1EA2\u1EA4\u1EA6\u1EA8\u1EAA\u1EAC\u1EAE\u1EB0\u1EB2\u1EB4\u1EB6\u1EB8\u1EBA\u1EBC\u1EBE\u1EC0\u1EC2\u1EC4\u1EC6\u1EC8\u1ECA\u1ECC\u1ECE\u1ED0\u1ED2\u1ED4\u1ED6\u1ED8\u1EDA\u1EDC\u1EDE\u1EE0\u1EE2\u1EE4\u1EE6\u1EE8\u1EEA\u1EEC\u1EEE\u1EF0\u1EF2\u1EF4\u1EF6\u1EF8\u1EFA\u1EFC\u1EFEЁА-ЯӘӨҮҖҢҺΑ-ΩΆΈΊΌΏΉΎА-ЩЮЯІЇЄҐЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F\'"”“`‘´’‚,„»«「」『』()〔〕【】《》〈〉])|(?<=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F]),(?=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])|(?<=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])(?:-|–|—|--|---|——|~)(?=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])|(?<=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F0-9])[:<>=/](?=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])�token_match��url_match�
|
|
|
|
| 2 |
��A�
|
| 3 |
� ��A� �'��A�'�''��A�''�(*_*)��A�(*_*)�(-8��A�(-8�(-:��A�(-:�(-;��A�(-;�(-_-)��A�(-_-)�(._.)��A�(._.)�(:��A�(:�(;��A�(;�(=��A�(=�(>_<)��A�(>_<)�(^_^)��A�(^_^)�(o:��A�(o:�(¬_¬)��A�(¬_¬)�(ಠ_ಠ)��A�(ಠ_ಠ)�(╯°□°)╯︵┻━┻��A�(╯°□°)╯︵┻━┻�)-:��A�)-:�):��A�):�-_-��A�-_-�-__-��A�-__-�._.��A�._.�0.0��A�0.0�0.o��A�0.o�0_0��A�0_0�0_o��A�0_o�8)��A�8)�8-)��A�8-)�8-D��A�8-D�8D��A�8D�:'(��A�:'(�:')��A�:')�:'-(��A�:'-(�:'-)��A�:'-)�:(��A�:(�:((��A�:((�:(((��A�:(((�:()��A�:()�:)��A�:)�:))��A�:))�:)))��A�:)))�:*��A�:*�:-(��A�:-(�:-((��A�:-((�:-(((��A�:-(((�:-)��A�:-)�:-))��A�:-))�:-)))��A�:-)))�:-*��A�:-*�:-/��A�:-/�:-0��A�:-0�:-3��A�:-3�:->��A�:->�:-D��A�:-D�:-O��A�:-O�:-P��A�:-P�:-X��A�:-X�:-]��A�:-]�:-o��A�:-o�:-p��A�:-p�:-x��A�:-x�:-|��A�:-|�:-}��A�:-}�:/��A�:/�:0��A�:0�:1��A�:1�:3��A�:3�:>��A�:>�:D��A�:D�:O��A�:O�:P��A�:P�:X��A�:X�:]��A�:]�:o��A�:o�:o)��A�:o)�:p��A�:p�:x��A�:x�:|��A�:|�:}��A�:}�:’(��A�:’(�:’)��A�:’)�:’-(��A�:’-(�:’-)��A�:’-)�;)��A�;)�;-)��A�;-)�;-D��A�;-D�;D��A�;D�;_;��A�;_;�<.<��A�<.<�</3��A�</3�<3��A�<3�<33��A�<33�<333��A�<333�<space>��A�<space>�=(��A�=(�=)��A�=)�=/��A�=/�=3��A�=3�=D��A�=D�=[��A�=[�=]��A�=]�=|��A�=|�>.<��A�>.<�>.>��A�>.>�>:(��A�>:(�>:o��A�>:o�><(((*>��A�><(((*>�@_@��A�@_@�Adm.��A�Adm.�Art.��A�Art.�Av.��A�Av.�C++��A�C++�Cia.��A�Cia.�Dr.��A�Dr.�E.G.��A�E.G.�E.g.��A�E.g.�Fund.��A�Fund.�Gen.��A�Gen.�Gov.��A�Gov.�I.E.��A�I.E.�I.e.��A�I.e.�Inc.��A�Inc.�Jr.��A�Jr.�Ltd.��A�Ltd.�Mr.��A�Mr.�O.O��A�O.O�O.o��A�O.o�O_O��A�O_O�O_o��A�O_o�Ph.D.��A�Ph.D.�Rep.��A�Rep.�Rev.��A�Rev.�S/A��A�S/A�Sen.��A�Sen.�Sr.��A�Sr.�Sra.��A�Sra.�V.V��A�V.V�V_V��A�V_V�XD��A�XD�XDD��A�XDD�[-:��A�[-:�[:��A�[:�[=��A�[=�\")��A�\")�\n��A�\n�\t��A�\t�]=��A�]=�^_^��A�^_^�^__^��A�^__^�^___^��A�^___^�a.��A�a.�art.��A�art.�av.��A�av.�b.��A�b.�c.��A�c.�d.��A�d.�dom.��A�dom.�dr.��A�dr.�e.��A�e.�e.g.��A�e.g.�e/ou��A�e/ou�ed.��A�ed.�eng.��A�eng.�etc.��A�etc.�f.��A�f.�g.��A�g.�h.��A�h.�i.��A�i.�i.e.��A�i.e.�j.��A�j.�k.��A�k.�km/h��A�km/h�l.��A�l.�m.��A�m.�n.��A�n.�o.��A�o.�o.0��A�o.0�o.O��A�o.O�o.o��A�o.o�o_0��A�o_0�o_O��A�o_O�o_o��A�o_o�p.��A�p.�p.m.��A�p.m.�pag.��A�pag.�pág.��A�pág.�q.��A�q.�r.��A�r.�s.��A�s.�sr.��A�sr.�sra.��A�sra.�t.��A�t.�tel.��A�tel.�u.��A�u.�v.��A�v.�v.v��A�v.v�v_v��A�v_v�vs.��A�vs.�w.��A�w.�x.��A�x.�xD��A�xD�xDD��A�xDD�y.��A�y.�z.��A�z.� ��A� C� �¯\(ツ)/¯��A�¯\(ツ)/¯�°C.��A�°�A�C�A�.�°F.��A�°�A�F�A�.�°K.��A�°�A�K�A�.�°c.��A�°�A�c�A�.�°f.��A�°�A�f�A�.�°k.��A�°�A�k�A�.�ä.��A�ä.�ö.��A�ö.�ü.��A�ü.�ಠ_ಠ��A�ಠ_ಠ�ಠ︵ಠ��A�ಠ︵ಠ�—��A�—�’��A�’�’’��A�’’�faster_heuristics�
|
|
|
|
| 1 |
+
��prefix_search��^\w{1,3}\$|^§|^%|^=|^—|^–|^\+(?![0-9])|^…|^……|^,|^:|^;|^\!|^\?|^¿|^؟|^¡|^\(|^\)|^\[|^\]|^\{|^\}|^<|^>|^_|^#|^\*|^&|^。|^?|^!|^,|^、|^;|^:|^~|^·|^।|^،|^۔|^؛|^٪|^\.\.+|^…|^\'|^"|^”|^“|^`|^‘|^´|^’|^‚|^,|^„|^»|^«|^「|^」|^『|^』|^(|^)|^〔|^〕|^【|^】|^《|^》|^〈|^〉|^〈|^〉|^⟦|^⟧|^\$|^£|^€|^¥|^฿|^US\$|^C\$|^A\$|^₽|^﷼|^₴|^₠|^₡|^₢|^₣|^₤|^₥|^₦|^₧|^₨|^₩|^₪|^₫|^€|^₭|^₮|^₯|^₰|^₱|^₲|^₳|^₴|^₵|^₶|^₷|^₸|^₹|^₺|^₻|^₼|^₽|^₾|^₿|^[\u00A6\u00A9\u00AE\u00B0\u0482\u058D\u058E\u060E\u060F\u06DE\u06E9\u06FD\u06FE\u07F6\u09FA\u0B70\u0BF3-\u0BF8\u0BFA\u0C7F\u0D4F\u0D79\u0F01-\u0F03\u0F13\u0F15-\u0F17\u0F1A-\u0F1F\u0F34\u0F36\u0F38\u0FBE-\u0FC5\u0FC7-\u0FCC\u0FCE\u0FCF\u0FD5-\u0FD8\u109E\u109F\u1390-\u1399\u1940\u19DE-\u19FF\u1B61-\u1B6A\u1B74-\u1B7C\u2100\u2101\u2103-\u2106\u2108\u2109\u2114\u2116\u2117\u211E-\u2123\u2125\u2127\u2129\u212E\u213A\u213B\u214A\u214C\u214D\u214F\u218A\u218B\u2195-\u2199\u219C-\u219F\u21A1\u21A2\u21A4\u21A5\u21A7-\u21AD\u21AF-\u21CD\u21D0\u21D1\u21D3\u21D5-\u21F3\u2300-\u2307\u230C-\u231F\u2322-\u2328\u232B-\u237B\u237D-\u239A\u23B4-\u23DB\u23E2-\u2426\u2440-\u244A\u249C-\u24E9\u2500-\u25B6\u25B8-\u25C0\u25C2-\u25F7\u2600-\u266E\u2670-\u2767\u2794-\u27BF\u2800-\u28FF\u2B00-\u2B2F\u2B45\u2B46\u2B4D-\u2B73\u2B76-\u2B95\u2B98-\u2BC8\u2BCA-\u2BFE\u2CE5-\u2CEA\u2E80-\u2E99\u2E9B-\u2EF3\u2F00-\u2FD5\u2FF0-\u2FFB\u3004\u3012\u3013\u3020\u3036\u3037\u303E\u303F\u3190\u3191\u3196-\u319F\u31C0-\u31E3\u3200-\u321E\u322A-\u3247\u3250\u3260-\u327F\u328A-\u32B0\u32C0-\u32FE\u3300-\u33FF\u4DC0-\u4DFF\uA490-\uA4C6\uA828-\uA82B\uA836\uA837\uA839\uAA77-\uAA79\uFDFD\uFFE4\uFFE8\uFFED\uFFEE\uFFFC\uFFFD\U00010137-\U0001013F\U00010179-\U00010189\U0001018C-\U0001018E\U00010190-\U0001019B\U000101A0\U000101D0-\U000101FC\U00010877\U00010878\U00010AC8\U0001173F\U00016B3C-\U00016B3F\U00016B45\U0001BC9C\U0001D000-\U0001D0F5\U0001D100-\U0001D126\U0001D129-\U0001D164\U0001D16A-\U0001D16C\U0001D183\U0001D184\U0001D18C-\U0001D1A9\U0001D1AE-\U0001D1E8\U0001D200-\U0001D241\U0001D245\U0001D300-\U0001D356\U0001D800-\U0001D9FF\U0001DA37-\U0001DA3A\U0001DA6D-\U0001DA74\U0001DA76-\U0001DA83\U0001DA85\U0001DA86\U0001ECAC\U0001F000-\U0001F02B\U0001F030-\U0001F093\U0001F0A0-\U0001F0AE\U0001F0B1-\U0001F0BF\U0001F0C1-\U0001F0CF\U0001F0D1-\U0001F0F5\U0001F110-\U0001F16B\U0001F170-\U0001F1AC\U0001F1E6-\U0001F202\U0001F210-\U0001F23B\U0001F240-\U0001F248\U0001F250\U0001F251\U0001F260-\U0001F265\U0001F300-\U0001F3FA\U0001F400-\U0001F6D4\U0001F6E0-\U0001F6EC\U0001F6F0-\U0001F6F9\U0001F700-\U0001F773\U0001F780-\U0001F7D8\U0001F800-\U0001F80B\U0001F810-\U0001F847\U0001F850-\U0001F859\U0001F860-\U0001F887\U0001F890-\U0001F8AD\U0001F900-\U0001F90B\U0001F910-\U0001F93E\U0001F940-\U0001F970\U0001F973-\U0001F976\U0001F97A\U0001F97C-\U0001F9A2\U0001F9B0-\U0001F9B9\U0001F9C0-\U0001F9C2\U0001F9D0-\U0001F9FF\U0001FA60-\U0001FA6D]�suffix_search�2�…$|……$|,$|:$|;$|\!$|\?$|¿$|؟$|¡$|\($|\)$|\[$|\]$|\{$|\}$|<$|>$|_$|#$|\*$|&$|。$|?$|!$|,$|、$|;$|:$|~$|·$|।$|،$|۔$|؛$|٪$|\.\.+$|…$|\'$|"$|”$|“$|`$|‘$|´$|’$|‚$|,$|„$|»$|«$|「$|」$|『$|』$|($|)$|〔$|〕$|【$|】$|《$|》$|〈$|〉$|〈$|〉$|⟦$|⟧$|[\u00A6\u00A9\u00AE\u00B0\u0482\u058D\u058E\u060E\u060F\u06DE\u06E9\u06FD\u06FE\u07F6\u09FA\u0B70\u0BF3-\u0BF8\u0BFA\u0C7F\u0D4F\u0D79\u0F01-\u0F03\u0F13\u0F15-\u0F17\u0F1A-\u0F1F\u0F34\u0F36\u0F38\u0FBE-\u0FC5\u0FC7-\u0FCC\u0FCE\u0FCF\u0FD5-\u0FD8\u109E\u109F\u1390-\u1399\u1940\u19DE-\u19FF\u1B61-\u1B6A\u1B74-\u1B7C\u2100\u2101\u2103-\u2106\u2108\u2109\u2114\u2116\u2117\u211E-\u2123\u2125\u2127\u2129\u212E\u213A\u213B\u214A\u214C\u214D\u214F\u218A\u218B\u2195-\u2199\u219C-\u219F\u21A1\u21A2\u21A4\u21A5\u21A7-\u21AD\u21AF-\u21CD\u21D0\u21D1\u21D3\u21D5-\u21F3\u2300-\u2307\u230C-\u231F\u2322-\u2328\u232B-\u237B\u237D-\u239A\u23B4-\u23DB\u23E2-\u2426\u2440-\u244A\u249C-\u24E9\u2500-\u25B6\u25B8-\u25C0\u25C2-\u25F7\u2600-\u266E\u2670-\u2767\u2794-\u27BF\u2800-\u28FF\u2B00-\u2B2F\u2B45\u2B46\u2B4D-\u2B73\u2B76-\u2B95\u2B98-\u2BC8\u2BCA-\u2BFE\u2CE5-\u2CEA\u2E80-\u2E99\u2E9B-\u2EF3\u2F00-\u2FD5\u2FF0-\u2FFB\u3004\u3012\u3013\u3020\u3036\u3037\u303E\u303F\u3190\u3191\u3196-\u319F\u31C0-\u31E3\u3200-\u321E\u322A-\u3247\u3250\u3260-\u327F\u328A-\u32B0\u32C0-\u32FE\u3300-\u33FF\u4DC0-\u4DFF\uA490-\uA4C6\uA828-\uA82B\uA836\uA837\uA839\uAA77-\uAA79\uFDFD\uFFE4\uFFE8\uFFED\uFFEE\uFFFC\uFFFD\U00010137-\U0001013F\U00010179-\U00010189\U0001018C-\U0001018E\U00010190-\U0001019B\U000101A0\U000101D0-\U000101FC\U00010877\U00010878\U00010AC8\U0001173F\U00016B3C-\U00016B3F\U00016B45\U0001BC9C\U0001D000-\U0001D0F5\U0001D100-\U0001D126\U0001D129-\U0001D164\U0001D16A-\U0001D16C\U0001D183\U0001D184\U0001D18C-\U0001D1A9\U0001D1AE-\U0001D1E8\U0001D200-\U0001D241\U0001D245\U0001D300-\U0001D356\U0001D800-\U0001D9FF\U0001DA37-\U0001DA3A\U0001DA6D-\U0001DA74\U0001DA76-\U0001DA83\U0001DA85\U0001DA86\U0001ECAC\U0001F000-\U0001F02B\U0001F030-\U0001F093\U0001F0A0-\U0001F0AE\U0001F0B1-\U0001F0BF\U0001F0C1-\U0001F0CF\U0001F0D1-\U0001F0F5\U0001F110-\U0001F16B\U0001F170-\U0001F1AC\U0001F1E6-\U0001F202\U0001F210-\U0001F23B\U0001F240-\U0001F248\U0001F250\U0001F251\U0001F260-\U0001F265\U0001F300-\U0001F3FA\U0001F400-\U0001F6D4\U0001F6E0-\U0001F6EC\U0001F6F0-\U0001F6F9\U0001F700-\U0001F773\U0001F780-\U0001F7D8\U0001F800-\U0001F80B\U0001F810-\U0001F847\U0001F850-\U0001F859\U0001F860-\U0001F887\U0001F890-\U0001F8AD\U0001F900-\U0001F90B\U0001F910-\U0001F93E\U0001F940-\U0001F970\U0001F973-\U0001F976\U0001F97A\U0001F97C-\U0001F9A2\U0001F9B0-\U0001F9B9\U0001F9C0-\U0001F9C2\U0001F9D0-\U0001F9FF\U0001FA60-\U0001FA6D]$|'s$|'S$|’s$|’S$|—$|–$|(?<=[0-9])\+$|(?<=°[FfCcKk])\.$|(?<=[0-9])(?:\$|£|€|¥|฿|US\$|C\$|A\$|₽|﷼|₴|₠|₡|₢|₣|₤|₥|₦|₧|₨|₩|₪|₫|€|₭|₮|₯|₰|₱|₲|₳|₴|₵|₶|₷|₸|₹|₺|₻|₼|₽|₾|₿)$|(?<=[0-9])(?:km|km²|km³|m|m²|m³|dm|dm²|dm³|cm|cm²|cm³|mm|mm²|mm³|ha|µm|nm|yd|in|ft|kg|g|mg|µg|t|lb|oz|m/s|km/h|kmh|mph|hPa|Pa|mbar|mb|MB|kb|KB|gb|GB|tb|TB|T|G|M|K|%|км|км²|км³|м|м²|м³|дм|дм²|дм³|см|см²|см³|мм|мм²|мм³|нм|кг|г|мг|м/с|км/ч|кПа|Па|мбар|Кб|КБ|кб|Мб|МБ|мб|Гб|ГБ|гб|Тб|ТБ|тбكم|كم²|كم³|م|م²|م³|سم|سم²|سم³|مم|مم²|مم³|كم|غرام|جرام|جم|كغ|ملغ|كوب|اكواب)$|(?<=[0-9a-z\uFF41-\uFF5A\u00DF-\u00F6\u00F8-\u00FF\u0101\u0103\u0105\u0107\u0109\u010B\u010D\u010F\u0111\u0113\u0115\u0117\u0119\u011B\u011D\u011F\u0121\u0123\u0125\u0127\u0129\u012B\u012D\u012F\u0131\u0133\u0135\u0137\u0138\u013A\u013C\u013E\u0140\u0142\u0144\u0146\u0148\u0149\u014B\u014D\u014F\u0151\u0153\u0155\u0157\u0159\u015B\u015D\u015F\u0161\u0163\u0165\u0167\u0169\u016B\u016D\u016F\u0171\u0173\u0175\u0177\u017A\u017C\u017E\u017F\u0180\u0183\u0185\u0188\u018C\u018D\u0192\u0195\u0199-\u019B\u019E\u01A1\u01A3\u01A5\u01A8\u01AA\u01AB\u01AD\u01B0\u01B4\u01B6\u01B9\u01BA\u01BD-\u01BF\u01C6\u01C9\u01CC\u01CE\u01D0\u01D2\u01D4\u01D6\u01D8\u01DA\u01DC\u01DD\u01DF\u01E1\u01E3\u01E5\u01E7\u01E9\u01EB\u01ED\u01EF\u01F0\u01F3\u01F5\u01F9\u01FB\u01FD\u01FF\u0201\u0203\u0205\u0207\u0209\u020B\u020D\u020F\u0211\u0213\u0215\u0217\u0219\u021B\u021D\u021F\u0221\u0223\u0225\u0227\u0229\u022B\u022D\u022F\u0231\u0233-\u0239\u023C\u023F\u0240\u0242\u0247\u0249\u024B\u024D\u024F\u2C61\u2C65\u2C66\u2C68\u2C6A\u2C6C\u2C71\u2C73\u2C74\u2C76-\u2C7B\uA723\uA725\uA727\uA729\uA72B\uA72D\uA72F-\uA731\uA733\uA735\uA737\uA739\uA73B\uA73D\uA73F\uA741\uA743\uA745\uA747\uA749\uA74B\uA74D\uA74F\uA751\uA753\uA755\uA757\uA759\uA75B\uA75D\uA75F\uA761\uA763\uA765\uA767\uA769\uA76B\uA76D\uA76F\uA771-\uA778\uA77A\uA77C\uA77F\uA781\uA783\uA785\uA787\uA78C\uA78E\uA791\uA793-\uA795\uA797\uA799\uA79B\uA79D\uA79F\uA7A1\uA7A3\uA7A5\uA7A7\uA7A9\uA7AF\uA7B5\uA7B7\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E01\u1E03\u1E05\u1E07\u1E09\u1E0B\u1E0D\u1E0F\u1E11\u1E13\u1E15\u1E17\u1E19\u1E1B\u1E1D\u1E1F\u1E21\u1E23\u1E25\u1E27\u1E29\u1E2B\u1E2D\u1E2F\u1E31\u1E33\u1E35\u1E37\u1E39\u1E3B\u1E3D\u1E3F\u1E41\u1E43\u1E45\u1E47\u1E49\u1E4B\u1E4D\u1E4F\u1E51\u1E53\u1E55\u1E57\u1E59\u1E5B\u1E5D\u1E5F\u1E61\u1E63\u1E65\u1E67\u1E69\u1E6B\u1E6D\u1E6F\u1E71\u1E73\u1E75\u1E77\u1E79\u1E7B\u1E7D\u1E7F\u1E81\u1E83\u1E85\u1E87\u1E89\u1E8B\u1E8D\u1E8F\u1E91\u1E93\u1E95-\u1E9D\u1E9F\u1EA1\u1EA3\u1EA5\u1EA7\u1EA9\u1EAB\u1EAD\u1EAF\u1EB1\u1EB3\u1EB5\u1EB7\u1EB9\u1EBB\u1EBD\u1EBF\u1EC1\u1EC3\u1EC5\u1EC7\u1EC9\u1ECB\u1ECD\u1ECF\u1ED1\u1ED3\u1ED5\u1ED7\u1ED9\u1EDB\u1EDD\u1EDF\u1EE1\u1EE3\u1EE5\u1EE7\u1EE9\u1EEB\u1EED\u1EEF\u1EF1\u1EF3\u1EF5\u1EF7\u1EF9\u1EFB\u1EFD\u1EFFёа-яәөүҗңһα-ωάέίόώήύа-щюяіїєґѓѕјљњќѐѝ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F%²\-\+…|……|,|:|;|\!|\?|¿|؟|¡|\(|\)|\[|\]|\{|\}|<|>|_|#|\*|&|。|?|!|,|、|;|:|~|·|।|،|۔|؛|٪(?:\'"”“`‘´’‚,„»«「」『』()〔〕【】《》〈〉〈〉⟦⟧)])\.$|(?<=[A-Z\uFF21-\uFF3A\u00C0-\u00D6\u00D8-\u00DE\u0100\u0102\u0104\u0106\u0108\u010A\u010C\u010E\u0110\u0112\u0114\u0116\u0118\u011A\u011C\u011E\u0120\u0122\u0124\u0126\u0128\u012A\u012C\u012E\u0130\u0132\u0134\u0136\u0139\u013B\u013D\u013F\u0141\u0143\u0145\u0147\u014A\u014C\u014E\u0150\u0152\u0154\u0156\u0158\u015A\u015C\u015E\u0160\u0162\u0164\u0166\u0168\u016A\u016C\u016E\u0170\u0172\u0174\u0176\u0178\u0179\u017B\u017D\u0181\u0182\u0184\u0186\u0187\u0189-\u018B\u018E-\u0191\u0193\u0194\u0196-\u0198\u019C\u019D\u019F\u01A0\u01A2\u01A4\u01A6\u01A7\u01A9\u01AC\u01AE\u01AF\u01B1-\u01B3\u01B5\u01B7\u01B8\u01BC\u01C4\u01C7\u01CA\u01CD\u01CF\u01D1\u01D3\u01D5\u01D7\u01D9\u01DB\u01DE\u01E0\u01E2\u01E4\u01E6\u01E8\u01EA\u01EC\u01EE\u01F1\u01F4\u01F6-\u01F8\u01FA\u01FC\u01FE\u0200\u0202\u0204\u0206\u0208\u020A\u020C\u020E\u0210\u0212\u0214\u0216\u0218\u021A\u021C\u021E\u0220\u0222\u0224\u0226\u0228\u022A\u022C\u022E\u0230\u0232\u023A\u023B\u023D\u023E\u0241\u0243-\u0246\u0248\u024A\u024C\u024E\u2C60\u2C62-\u2C64\u2C67\u2C69\u2C6B\u2C6D-\u2C70\u2C72\u2C75\u2C7E\u2C7F\uA722\uA724\uA726\uA728\uA72A\uA72C\uA72E\uA732\uA734\uA736\uA738\uA73A\uA73C\uA73E\uA740\uA742\uA744\uA746\uA748\uA74A\uA74C\uA74E\uA750\uA752\uA754\uA756\uA758\uA75A\uA75C\uA75E\uA760\uA762\uA764\uA766\uA768\uA76A\uA76C\uA76E\uA779\uA77B\uA77D\uA77E\uA780\uA782\uA784\uA786\uA78B\uA78D\uA790\uA792\uA796\uA798\uA79A\uA79C\uA79E\uA7A0\uA7A2\uA7A4\uA7A6\uA7A8\uA7AA-\uA7AE\uA7B0-\uA7B4\uA7B6\uA7B8\u1E00\u1E02\u1E04\u1E06\u1E08\u1E0A\u1E0C\u1E0E\u1E10\u1E12\u1E14\u1E16\u1E18\u1E1A\u1E1C\u1E1E\u1E20\u1E22\u1E24\u1E26\u1E28\u1E2A\u1E2C\u1E2E\u1E30\u1E32\u1E34\u1E36\u1E38\u1E3A\u1E3C\u1E3E\u1E40\u1E42\u1E44\u1E46\u1E48\u1E4A\u1E4C\u1E4E\u1E50\u1E52\u1E54\u1E56\u1E58\u1E5A\u1E5C\u1E5E\u1E60\u1E62\u1E64\u1E66\u1E68\u1E6A\u1E6C\u1E6E\u1E70\u1E72\u1E74\u1E76\u1E78\u1E7A\u1E7C\u1E7E\u1E80\u1E82\u1E84\u1E86\u1E88\u1E8A\u1E8C\u1E8E\u1E90\u1E92\u1E94\u1E9E\u1EA0\u1EA2\u1EA4\u1EA6\u1EA8\u1EAA\u1EAC\u1EAE\u1EB0\u1EB2\u1EB4\u1EB6\u1EB8\u1EBA\u1EBC\u1EBE\u1EC0\u1EC2\u1EC4\u1EC6\u1EC8\u1ECA\u1ECC\u1ECE\u1ED0\u1ED2\u1ED4\u1ED6\u1ED8\u1EDA\u1EDC\u1EDE\u1EE0\u1EE2\u1EE4\u1EE6\u1EE8\u1EEA\u1EEC\u1EEE\u1EF0\u1EF2\u1EF4\u1EF6\u1EF8\u1EFA\u1EFC\u1EFEЁА-ЯӘӨҮҖҢҺΑ-ΩΆΈΊΌΏΉΎА-ЩЮЯІЇЄҐЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F][A-Z\uFF21-\uFF3A\u00C0-\u00D6\u00D8-\u00DE\u0100\u0102\u0104\u0106\u0108\u010A\u010C\u010E\u0110\u0112\u0114\u0116\u0118\u011A\u011C\u011E\u0120\u0122\u0124\u0126\u0128\u012A\u012C\u012E\u0130\u0132\u0134\u0136\u0139\u013B\u013D\u013F\u0141\u0143\u0145\u0147\u014A\u014C\u014E\u0150\u0152\u0154\u0156\u0158\u015A\u015C\u015E\u0160\u0162\u0164\u0166\u0168\u016A\u016C\u016E\u0170\u0172\u0174\u0176\u0178\u0179\u017B\u017D\u0181\u0182\u0184\u0186\u0187\u0189-\u018B\u018E-\u0191\u0193\u0194\u0196-\u0198\u019C\u019D\u019F\u01A0\u01A2\u01A4\u01A6\u01A7\u01A9\u01AC\u01AE\u01AF\u01B1-\u01B3\u01B5\u01B7\u01B8\u01BC\u01C4\u01C7\u01CA\u01CD\u01CF\u01D1\u01D3\u01D5\u01D7\u01D9\u01DB\u01DE\u01E0\u01E2\u01E4\u01E6\u01E8\u01EA\u01EC\u01EE\u01F1\u01F4\u01F6-\u01F8\u01FA\u01FC\u01FE\u0200\u0202\u0204\u0206\u0208\u020A\u020C\u020E\u0210\u0212\u0214\u0216\u0218\u021A\u021C\u021E\u0220\u0222\u0224\u0226\u0228\u022A\u022C\u022E\u0230\u0232\u023A\u023B\u023D\u023E\u0241\u0243-\u0246\u0248\u024A\u024C\u024E\u2C60\u2C62-\u2C64\u2C67\u2C69\u2C6B\u2C6D-\u2C70\u2C72\u2C75\u2C7E\u2C7F\uA722\uA724\uA726\uA728\uA72A\uA72C\uA72E\uA732\uA734\uA736\uA738\uA73A\uA73C\uA73E\uA740\uA742\uA744\uA746\uA748\uA74A\uA74C\uA74E\uA750\uA752\uA754\uA756\uA758\uA75A\uA75C\uA75E\uA760\uA762\uA764\uA766\uA768\uA76A\uA76C\uA76E\uA779\uA77B\uA77D\uA77E\uA780\uA782\uA784\uA786\uA78B\uA78D\uA790\uA792\uA796\uA798\uA79A\uA79C\uA79E\uA7A0\uA7A2\uA7A4\uA7A6\uA7A8\uA7AA-\uA7AE\uA7B0-\uA7B4\uA7B6\uA7B8\u1E00\u1E02\u1E04\u1E06\u1E08\u1E0A\u1E0C\u1E0E\u1E10\u1E12\u1E14\u1E16\u1E18\u1E1A\u1E1C\u1E1E\u1E20\u1E22\u1E24\u1E26\u1E28\u1E2A\u1E2C\u1E2E\u1E30\u1E32\u1E34\u1E36\u1E38\u1E3A\u1E3C\u1E3E\u1E40\u1E42\u1E44\u1E46\u1E48\u1E4A\u1E4C\u1E4E\u1E50\u1E52\u1E54\u1E56\u1E58\u1E5A\u1E5C\u1E5E\u1E60\u1E62\u1E64\u1E66\u1E68\u1E6A\u1E6C\u1E6E\u1E70\u1E72\u1E74\u1E76\u1E78\u1E7A\u1E7C\u1E7E\u1E80\u1E82\u1E84\u1E86\u1E88\u1E8A\u1E8C\u1E8E\u1E90\u1E92\u1E94\u1E9E\u1EA0\u1EA2\u1EA4\u1EA6\u1EA8\u1EAA\u1EAC\u1EAE\u1EB0\u1EB2\u1EB4\u1EB6\u1EB8\u1EBA\u1EBC\u1EBE\u1EC0\u1EC2\u1EC4\u1EC6\u1EC8\u1ECA\u1ECC\u1ECE\u1ED0\u1ED2\u1ED4\u1ED6\u1ED8\u1EDA\u1EDC\u1EDE\u1EE0\u1EE2\u1EE4\u1EE6\u1EE8\u1EEA\u1EEC\u1EEE\u1EF0\u1EF2\u1EF4\u1EF6\u1EF8\u1EFA\u1EFC\u1EFEЁА-ЯӘӨҮҖҢҺΑ-ΩΆΈΊΌΏΉΎА-ЩЮЯІЇЄҐЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])\.$�infix_finditer�?
|
| 2 |
+
(\w+-\w+(-\w+)*)|\.\.+|…|[\u00A6\u00A9\u00AE\u00B0\u0482\u058D\u058E\u060E\u060F\u06DE\u06E9\u06FD\u06FE\u07F6\u09FA\u0B70\u0BF3-\u0BF8\u0BFA\u0C7F\u0D4F\u0D79\u0F01-\u0F03\u0F13\u0F15-\u0F17\u0F1A-\u0F1F\u0F34\u0F36\u0F38\u0FBE-\u0FC5\u0FC7-\u0FCC\u0FCE\u0FCF\u0FD5-\u0FD8\u109E\u109F\u1390-\u1399\u1940\u19DE-\u19FF\u1B61-\u1B6A\u1B74-\u1B7C\u2100\u2101\u2103-\u2106\u2108\u2109\u2114\u2116\u2117\u211E-\u2123\u2125\u2127\u2129\u212E\u213A\u213B\u214A\u214C\u214D\u214F\u218A\u218B\u2195-\u2199\u219C-\u219F\u21A1\u21A2\u21A4\u21A5\u21A7-\u21AD\u21AF-\u21CD\u21D0\u21D1\u21D3\u21D5-\u21F3\u2300-\u2307\u230C-\u231F\u2322-\u2328\u232B-\u237B\u237D-\u239A\u23B4-\u23DB\u23E2-\u2426\u2440-\u244A\u249C-\u24E9\u2500-\u25B6\u25B8-\u25C0\u25C2-\u25F7\u2600-\u266E\u2670-\u2767\u2794-\u27BF\u2800-\u28FF\u2B00-\u2B2F\u2B45\u2B46\u2B4D-\u2B73\u2B76-\u2B95\u2B98-\u2BC8\u2BCA-\u2BFE\u2CE5-\u2CEA\u2E80-\u2E99\u2E9B-\u2EF3\u2F00-\u2FD5\u2FF0-\u2FFB\u3004\u3012\u3013\u3020\u3036\u3037\u303E\u303F\u3190\u3191\u3196-\u319F\u31C0-\u31E3\u3200-\u321E\u322A-\u3247\u3250\u3260-\u327F\u328A-\u32B0\u32C0-\u32FE\u3300-\u33FF\u4DC0-\u4DFF\uA490-\uA4C6\uA828-\uA82B\uA836\uA837\uA839\uAA77-\uAA79\uFDFD\uFFE4\uFFE8\uFFED\uFFEE\uFFFC\uFFFD\U00010137-\U0001013F\U00010179-\U00010189\U0001018C-\U0001018E\U00010190-\U0001019B\U000101A0\U000101D0-\U000101FC\U00010877\U00010878\U00010AC8\U0001173F\U00016B3C-\U00016B3F\U00016B45\U0001BC9C\U0001D000-\U0001D0F5\U0001D100-\U0001D126\U0001D129-\U0001D164\U0001D16A-\U0001D16C\U0001D183\U0001D184\U0001D18C-\U0001D1A9\U0001D1AE-\U0001D1E8\U0001D200-\U0001D241\U0001D245\U0001D300-\U0001D356\U0001D800-\U0001D9FF\U0001DA37-\U0001DA3A\U0001DA6D-\U0001DA74\U0001DA76-\U0001DA83\U0001DA85\U0001DA86\U0001ECAC\U0001F000-\U0001F02B\U0001F030-\U0001F093\U0001F0A0-\U0001F0AE\U0001F0B1-\U0001F0BF\U0001F0C1-\U0001F0CF\U0001F0D1-\U0001F0F5\U0001F110-\U0001F16B\U0001F170-\U0001F1AC\U0001F1E6-\U0001F202\U0001F210-\U0001F23B\U0001F240-\U0001F248\U0001F250\U0001F251\U0001F260-\U0001F265\U0001F300-\U0001F3FA\U0001F400-\U0001F6D4\U0001F6E0-\U0001F6EC\U0001F6F0-\U0001F6F9\U0001F700-\U0001F773\U0001F780-\U0001F7D8\U0001F800-\U0001F80B\U0001F810-\U0001F847\U0001F850-\U0001F859\U0001F860-\U0001F887\U0001F890-\U0001F8AD\U0001F900-\U0001F90B\U0001F910-\U0001F93E\U0001F940-\U0001F970\U0001F973-\U0001F976\U0001F97A\U0001F97C-\U0001F9A2\U0001F9B0-\U0001F9B9\U0001F9C0-\U0001F9C2\U0001F9D0-\U0001F9FF\U0001FA60-\U0001FA6D]|(?<=[0-9])[+\-\*^](?=[0-9-])|(?<=[a-z\uFF41-\uFF5A\u00DF-\u00F6\u00F8-\u00FF\u0101\u0103\u0105\u0107\u0109\u010B\u010D\u010F\u0111\u0113\u0115\u0117\u0119\u011B\u011D\u011F\u0121\u0123\u0125\u0127\u0129\u012B\u012D\u012F\u0131\u0133\u0135\u0137\u0138\u013A\u013C\u013E\u0140\u0142\u0144\u0146\u0148\u0149\u014B\u014D\u014F\u0151\u0153\u0155\u0157\u0159\u015B\u015D\u015F\u0161\u0163\u0165\u0167\u0169\u016B\u016D\u016F\u0171\u0173\u0175\u0177\u017A\u017C\u017E\u017F\u0180\u0183\u0185\u0188\u018C\u018D\u0192\u0195\u0199-\u019B\u019E\u01A1\u01A3\u01A5\u01A8\u01AA\u01AB\u01AD\u01B0\u01B4\u01B6\u01B9\u01BA\u01BD-\u01BF\u01C6\u01C9\u01CC\u01CE\u01D0\u01D2\u01D4\u01D6\u01D8\u01DA\u01DC\u01DD\u01DF\u01E1\u01E3\u01E5\u01E7\u01E9\u01EB\u01ED\u01EF\u01F0\u01F3\u01F5\u01F9\u01FB\u01FD\u01FF\u0201\u0203\u0205\u0207\u0209\u020B\u020D\u020F\u0211\u0213\u0215\u0217\u0219\u021B\u021D\u021F\u0221\u0223\u0225\u0227\u0229\u022B\u022D\u022F\u0231\u0233-\u0239\u023C\u023F\u0240\u0242\u0247\u0249\u024B\u024D\u024F\u2C61\u2C65\u2C66\u2C68\u2C6A\u2C6C\u2C71\u2C73\u2C74\u2C76-\u2C7B\uA723\uA725\uA727\uA729\uA72B\uA72D\uA72F-\uA731\uA733\uA735\uA737\uA739\uA73B\uA73D\uA73F\uA741\uA743\uA745\uA747\uA749\uA74B\uA74D\uA74F\uA751\uA753\uA755\uA757\uA759\uA75B\uA75D\uA75F\uA761\uA763\uA765\uA767\uA769\uA76B\uA76D\uA76F\uA771-\uA778\uA77A\uA77C\uA77F\uA781\uA783\uA785\uA787\uA78C\uA78E\uA791\uA793-\uA795\uA797\uA799\uA79B\uA79D\uA79F\uA7A1\uA7A3\uA7A5\uA7A7\uA7A9\uA7AF\uA7B5\uA7B7\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E01\u1E03\u1E05\u1E07\u1E09\u1E0B\u1E0D\u1E0F\u1E11\u1E13\u1E15\u1E17\u1E19\u1E1B\u1E1D\u1E1F\u1E21\u1E23\u1E25\u1E27\u1E29\u1E2B\u1E2D\u1E2F\u1E31\u1E33\u1E35\u1E37\u1E39\u1E3B\u1E3D\u1E3F\u1E41\u1E43\u1E45\u1E47\u1E49\u1E4B\u1E4D\u1E4F\u1E51\u1E53\u1E55\u1E57\u1E59\u1E5B\u1E5D\u1E5F\u1E61\u1E63\u1E65\u1E67\u1E69\u1E6B\u1E6D\u1E6F\u1E71\u1E73\u1E75\u1E77\u1E79\u1E7B\u1E7D\u1E7F\u1E81\u1E83\u1E85\u1E87\u1E89\u1E8B\u1E8D\u1E8F\u1E91\u1E93\u1E95-\u1E9D\u1E9F\u1EA1\u1EA3\u1EA5\u1EA7\u1EA9\u1EAB\u1EAD\u1EAF\u1EB1\u1EB3\u1EB5\u1EB7\u1EB9\u1EBB\u1EBD\u1EBF\u1EC1\u1EC3\u1EC5\u1EC7\u1EC9\u1ECB\u1ECD\u1ECF\u1ED1\u1ED3\u1ED5\u1ED7\u1ED9\u1EDB\u1EDD\u1EDF\u1EE1\u1EE3\u1EE5\u1EE7\u1EE9\u1EEB\u1EED\u1EEF\u1EF1\u1EF3\u1EF5\u1EF7\u1EF9\u1EFB\u1EFD\u1EFFёа-яәөүҗңһα-ωάέίόώήύа-щюяіїєґѓѕјљњќѐѝ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F\'"”“`‘´’‚,„»«「」『』()〔〕【】《》〈〉〈〉⟦⟧])\.(?=[A-Z\uFF21-\uFF3A\u00C0-\u00D6\u00D8-\u00DE\u0100\u0102\u0104\u0106\u0108\u010A\u010C\u010E\u0110\u0112\u0114\u0116\u0118\u011A\u011C\u011E\u0120\u0122\u0124\u0126\u0128\u012A\u012C\u012E\u0130\u0132\u0134\u0136\u0139\u013B\u013D\u013F\u0141\u0143\u0145\u0147\u014A\u014C\u014E\u0150\u0152\u0154\u0156\u0158\u015A\u015C\u015E\u0160\u0162\u0164\u0166\u0168\u016A\u016C\u016E\u0170\u0172\u0174\u0176\u0178\u0179\u017B\u017D\u0181\u0182\u0184\u0186\u0187\u0189-\u018B\u018E-\u0191\u0193\u0194\u0196-\u0198\u019C\u019D\u019F\u01A0\u01A2\u01A4\u01A6\u01A7\u01A9\u01AC\u01AE\u01AF\u01B1-\u01B3\u01B5\u01B7\u01B8\u01BC\u01C4\u01C7\u01CA\u01CD\u01CF\u01D1\u01D3\u01D5\u01D7\u01D9\u01DB\u01DE\u01E0\u01E2\u01E4\u01E6\u01E8\u01EA\u01EC\u01EE\u01F1\u01F4\u01F6-\u01F8\u01FA\u01FC\u01FE\u0200\u0202\u0204\u0206\u0208\u020A\u020C\u020E\u0210\u0212\u0214\u0216\u0218\u021A\u021C\u021E\u0220\u0222\u0224\u0226\u0228\u022A\u022C\u022E\u0230\u0232\u023A\u023B\u023D\u023E\u0241\u0243-\u0246\u0248\u024A\u024C\u024E\u2C60\u2C62-\u2C64\u2C67\u2C69\u2C6B\u2C6D-\u2C70\u2C72\u2C75\u2C7E\u2C7F\uA722\uA724\uA726\uA728\uA72A\uA72C\uA72E\uA732\uA734\uA736\uA738\uA73A\uA73C\uA73E\uA740\uA742\uA744\uA746\uA748\uA74A\uA74C\uA74E\uA750\uA752\uA754\uA756\uA758\uA75A\uA75C\uA75E\uA760\uA762\uA764\uA766\uA768\uA76A\uA76C\uA76E\uA779\uA77B\uA77D\uA77E\uA780\uA782\uA784\uA786\uA78B\uA78D\uA790\uA792\uA796\uA798\uA79A\uA79C\uA79E\uA7A0\uA7A2\uA7A4\uA7A6\uA7A8\uA7AA-\uA7AE\uA7B0-\uA7B4\uA7B6\uA7B8\u1E00\u1E02\u1E04\u1E06\u1E08\u1E0A\u1E0C\u1E0E\u1E10\u1E12\u1E14\u1E16\u1E18\u1E1A\u1E1C\u1E1E\u1E20\u1E22\u1E24\u1E26\u1E28\u1E2A\u1E2C\u1E2E\u1E30\u1E32\u1E34\u1E36\u1E38\u1E3A\u1E3C\u1E3E\u1E40\u1E42\u1E44\u1E46\u1E48\u1E4A\u1E4C\u1E4E\u1E50\u1E52\u1E54\u1E56\u1E58\u1E5A\u1E5C\u1E5E\u1E60\u1E62\u1E64\u1E66\u1E68\u1E6A\u1E6C\u1E6E\u1E70\u1E72\u1E74\u1E76\u1E78\u1E7A\u1E7C\u1E7E\u1E80\u1E82\u1E84\u1E86\u1E88\u1E8A\u1E8C\u1E8E\u1E90\u1E92\u1E94\u1E9E\u1EA0\u1EA2\u1EA4\u1EA6\u1EA8\u1EAA\u1EAC\u1EAE\u1EB0\u1EB2\u1EB4\u1EB6\u1EB8\u1EBA\u1EBC\u1EBE\u1EC0\u1EC2\u1EC4\u1EC6\u1EC8\u1ECA\u1ECC\u1ECE\u1ED0\u1ED2\u1ED4\u1ED6\u1ED8\u1EDA\u1EDC\u1EDE\u1EE0\u1EE2\u1EE4\u1EE6\u1EE8\u1EEA\u1EEC\u1EEE\u1EF0\u1EF2\u1EF4\u1EF6\u1EF8\u1EFA\u1EFC\u1EFEЁА-ЯӘӨҮҖҢҺΑ-ΩΆΈΊΌΏΉΎА-ЩЮЯІЇЄҐЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F\'"”“`‘´’‚,„»«「」『』()〔〕【】《》〈〉〈〉⟦⟧])|(?<=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F]),(?=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])|(?<=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])(?:-|–|—|--|---|——|~)(?=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])|(?<=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F0-9])[:<>=/](?=[A-Za-z\uFF21-\uFF3A\uFF41-\uFF5A\u00C0-\u00D6\u00D8-\u00F6\u00F8-\u00FF\u0100-\u017F\u0180-\u01BF\u01C4-\u024F\u2C60-\u2C7B\u2C7E\u2C7F\uA722-\uA76F\uA771-\uA787\uA78B-\uA78E\uA790-\uA7B9\uA7FA\uAB30-\uAB5A\uAB60-\uAB64\u0250-\u02AF\u1D00-\u1D25\u1D6B-\u1D77\u1D79-\u1D9A\u1E00-\u1EFFёа-яЁА-ЯәөүҗңһӘӨҮҖҢҺα-ωάέίόώήύΑ-ΩΆΈΊΌΏΉΎа-щюяіїєґА-ЩЮЯІЇЄҐѓѕјљњќѐѝЃЅЈЉЊЌЀЍ\u1200-\u137F\u0980-\u09FF\u0591-\u05F4\uFB1D-\uFB4F\u0620-\u064A\u066E-\u06D5\u06E5-\u06FF\u0750-\u077F\u08A0-\u08BD\uFB50-\uFBB1\uFBD3-\uFD3D\uFD50-\uFDC7\uFDF0-\uFDFB\uFE70-\uFEFC\U0001EE00-\U0001EEBB\u0D80-\u0DFF\u0900-\u097F\u0C80-\u0CFF\u0B80-\u0BFF\u0C00-\u0C7F\uAC00-\uD7AF\u1100-\u11FF\u3040-\u309F\u30A0-\u30FFー\u4E00-\u62FF\u6300-\u77FF\u7800-\u8CFF\u8D00-\u9FFF\u3400-\u4DBF\U00020000-\U000215FF\U00021600-\U000230FF\U00023100-\U000245FF\U00024600-\U000260FF\U00026100-\U000275FF\U00027600-\U000290FF\U00029100-\U0002A6DF\U0002A700-\U0002B73F\U0002B740-\U0002B81F\U0002B820-\U0002CEAF\U0002CEB0-\U0002EBEF\u2E80-\u2EFF\u2F00-\u2FDF\u2FF0-\u2FFF\u3000-\u303F\u31C0-\u31EF\u3200-\u32FF\u3300-\u33FF\uF900-\uFAFF\uFE30-\uFE4F\U0001F200-\U0001F2FF\U0002F800-\U0002FA1F])�token_match��url_match�
|
| 3 |
��A�
|
| 4 |
� ��A� �'��A�'�''��A�''�(*_*)��A�(*_*)�(-8��A�(-8�(-:��A�(-:�(-;��A�(-;�(-_-)��A�(-_-)�(._.)��A�(._.)�(:��A�(:�(;��A�(;�(=��A�(=�(>_<)��A�(>_<)�(^_^)��A�(^_^)�(o:��A�(o:�(¬_¬)��A�(¬_¬)�(ಠ_ಠ)��A�(ಠ_ಠ)�(╯°□°)╯︵┻━┻��A�(╯°□°)╯︵┻━┻�)-:��A�)-:�):��A�):�-_-��A�-_-�-__-��A�-__-�._.��A�._.�0.0��A�0.0�0.o��A�0.o�0_0��A�0_0�0_o��A�0_o�8)��A�8)�8-)��A�8-)�8-D��A�8-D�8D��A�8D�:'(��A�:'(�:')��A�:')�:'-(��A�:'-(�:'-)��A�:'-)�:(��A�:(�:((��A�:((�:(((��A�:(((�:()��A�:()�:)��A�:)�:))��A�:))�:)))��A�:)))�:*��A�:*�:-(��A�:-(�:-((��A�:-((�:-(((��A�:-(((�:-)��A�:-)�:-))��A�:-))�:-)))��A�:-)))�:-*��A�:-*�:-/��A�:-/�:-0��A�:-0�:-3��A�:-3�:->��A�:->�:-D��A�:-D�:-O��A�:-O�:-P��A�:-P�:-X��A�:-X�:-]��A�:-]�:-o��A�:-o�:-p��A�:-p�:-x��A�:-x�:-|��A�:-|�:-}��A�:-}�:/��A�:/�:0��A�:0�:1��A�:1�:3��A�:3�:>��A�:>�:D��A�:D�:O��A�:O�:P��A�:P�:X��A�:X�:]��A�:]�:o��A�:o�:o)��A�:o)�:p��A�:p�:x��A�:x�:|��A�:|�:}��A�:}�:’(��A�:’(�:’)��A�:’)�:’-(��A�:’-(�:’-)��A�:’-)�;)��A�;)�;-)��A�;-)�;-D��A�;-D�;D��A�;D�;_;��A�;_;�<.<��A�<.<�</3��A�</3�<3��A�<3�<33��A�<33�<333��A�<333�<space>��A�<space>�=(��A�=(�=)��A�=)�=/��A�=/�=3��A�=3�=D��A�=D�=[��A�=[�=]��A�=]�=|��A�=|�>.<��A�>.<�>.>��A�>.>�>:(��A�>:(�>:o��A�>:o�><(((*>��A�><(((*>�@_@��A�@_@�Adm.��A�Adm.�Art.��A�Art.�Av.��A�Av.�C++��A�C++�Cia.��A�Cia.�Dr.��A�Dr.�E.G.��A�E.G.�E.g.��A�E.g.�Fund.��A�Fund.�Gen.��A�Gen.�Gov.��A�Gov.�I.E.��A�I.E.�I.e.��A�I.e.�Inc.��A�Inc.�Jr.��A�Jr.�Ltd.��A�Ltd.�Mr.��A�Mr.�O.O��A�O.O�O.o��A�O.o�O_O��A�O_O�O_o��A�O_o�Ph.D.��A�Ph.D.�Rep.��A�Rep.�Rev.��A�Rev.�S/A��A�S/A�Sen.��A�Sen.�Sr.��A�Sr.�Sra.��A�Sra.�V.V��A�V.V�V_V��A�V_V�XD��A�XD�XDD��A�XDD�[-:��A�[-:�[:��A�[:�[=��A�[=�\")��A�\")�\n��A�\n�\t��A�\t�]=��A�]=�^_^��A�^_^�^__^��A�^__^�^___^��A�^___^�a.��A�a.�art.��A�art.�av.��A�av.�b.��A�b.�c.��A�c.�d.��A�d.�dom.��A�dom.�dr.��A�dr.�e.��A�e.�e.g.��A�e.g.�e/ou��A�e/ou�ed.��A�ed.�eng.��A�eng.�etc.��A�etc.�f.��A�f.�g.��A�g.�h.��A�h.�i.��A�i.�i.e.��A�i.e.�j.��A�j.�k.��A�k.�km/h��A�km/h�l.��A�l.�m.��A�m.�n.��A�n.�o.��A�o.�o.0��A�o.0�o.O��A�o.O�o.o��A�o.o�o_0��A�o_0�o_O��A�o_O�o_o��A�o_o�p.��A�p.�p.m.��A�p.m.�pag.��A�pag.�pág.��A�pág.�q.��A�q.�r.��A�r.�s.��A�s.�sr.��A�sr.�sra.��A�sra.�t.��A�t.�tel.��A�tel.�u.��A�u.�v.��A�v.�v.v��A�v.v�v_v��A�v_v�vs.��A�vs.�w.��A�w.�x.��A�x.�xD��A�xD�xDD��A�xDD�y.��A�y.�z.��A�z.� ��A� C� �¯\(ツ)/¯��A�¯\(ツ)/¯�°C.��A�°�A�C�A�.�°F.��A�°�A�F�A�.�°K.��A�°�A�K�A�.�°c.��A�°�A�c�A�.�°f.��A�°�A�f�A�.�°k.��A�°�A�k�A�.�ä.��A�ä.�ö.��A�ö.�ü.��A�ü.�ಠ_ಠ��A�ಠ_ಠ�ಠ︵ಠ��A�ಠ︵ಠ�—��A�—�’��A�’�’’��A�’’�faster_heuristics�
|
vocab/strings.json
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
-
size
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:bd37b00976bb0a36fa9b2d19fae890834099c4407e83cbd89a978bc581103cec
|
| 3 |
+
size 9851704
|