fr_core_news_sm / meta.json
adrianeboyd's picture
Update spaCy pipeline
a0ad81c
{
"lang":"fr",
"name":"core_news_sm",
"version":"3.7.0",
"description":"French pipeline optimized for CPU. Components: tok2vec, morphologizer, parser, senter, ner, attribute_ruler, lemmatizer.",
"author":"Explosion",
"email":"contact@explosion.ai",
"url":"https://explosion.ai",
"license":"LGPL-LR",
"spacy_version":">=3.7.0,<3.8.0",
"spacy_git_version":"6b4f77441",
"vectors":{
"width":0,
"vectors":0,
"keys":0,
"name":null
},
"labels":{
"tok2vec":[
],
"morphologizer":[
"POS=PROPN",
"Gender=Fem|Number=Sing|POS=DET|PronType=Dem",
"Gender=Fem|Number=Sing|POS=NOUN",
"Number=Plur|POS=PRON|Person=1",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"POS=SCONJ",
"POS=ADP",
"Definite=Def|Gender=Masc|Number=Sing|POS=DET|PronType=Art",
"NumType=Ord|POS=ADJ",
"Gender=Masc|Number=Sing|POS=NOUN",
"POS=PUNCT",
"Gender=Masc|Number=Sing|POS=PROPN",
"Number=Plur|POS=ADJ",
"Gender=Masc|Number=Plur|POS=NOUN",
"Definite=Ind|Gender=Fem|Number=Sing|POS=DET|PronType=Art",
"Number=Sing|POS=ADJ",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Imp|VerbForm=Fin",
"POS=ADV",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Definite=Def|Gender=Fem|Number=Sing|POS=DET|PronType=Art",
"Gender=Fem|Number=Sing|POS=PROPN",
"Definite=Def|Number=Sing|POS=DET|PronType=Art",
"NumType=Card|POS=NUM",
"Definite=Def|Number=Plur|POS=DET|PronType=Art",
"Gender=Masc|Number=Plur|POS=ADJ",
"POS=CCONJ",
"Gender=Fem|Number=Plur|POS=NOUN",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Gender=Fem|Number=Plur|POS=ADJ",
"POS=ADJ",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"POS=PRON|PronType=Rel",
"Number=Sing|POS=DET|Poss=Yes",
"Definite=Def|Gender=Masc|Number=Sing|POS=ADP|PronType=Art",
"Definite=Def|Number=Plur|POS=ADP|PronType=Art",
"Definite=Ind|Number=Plur|POS=DET|PronType=Art",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Masc|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"POS=VERB|VerbForm=Inf",
"Gender=Fem|Number=Sing|POS=ADJ",
"Gender=Masc|Number=Sing|POS=PRON|Person=3",
"Number=Plur|POS=DET",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=ADJ",
"Gender=Masc|Number=Sing|POS=DET|PronType=Dem",
"POS=ADV|PronType=Int",
"POS=VERB|Tense=Pres|VerbForm=Part",
"Gender=Fem|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Ind|Gender=Masc|Number=Sing|POS=DET|PronType=Art",
"Gender=Masc|POS=ADJ",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Fut|VerbForm=Fin",
"Number=Plur|POS=DET|Poss=Yes",
"POS=AUX|VerbForm=Inf",
"Gender=Masc|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Masc|POS=VERB|Tense=Past|VerbForm=Part",
"POS=ADV|Polarity=Neg",
"Definite=Ind|Number=Sing|POS=DET|PronType=Art",
"Gender=Fem|Number=Sing|POS=PRON|Person=3",
"POS=PRON|Person=3|Reflex=Yes",
"Gender=Masc|POS=NOUN",
"POS=AUX|Tense=Past|VerbForm=Part",
"POS=PRON|Person=3",
"Number=Plur|POS=NOUN",
"NumType=Ord|Number=Sing|POS=ADJ",
"POS=VERB|Tense=Past|VerbForm=Part",
"POS=AUX|Tense=Pres|VerbForm=Part",
"Gender=Masc|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Sing|POS=PRON|Person=3",
"Number=Sing|POS=NOUN",
"Gender=Masc|Number=Plur|POS=PRON|Person=3",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Imp|VerbForm=Fin",
"Gender=Fem|NumType=Ord|Number=Sing|POS=ADJ",
"Number=Plur|POS=PROPN",
"Number=Sing|POS=PROPN",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Plur|POS=PRON|PronType=Dem",
"Gender=Masc|Number=Sing|POS=DET",
"Gender=Fem|Number=Sing|POS=DET|Poss=Yes",
"Gender=Masc|POS=PRON",
"POS=NOUN",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON",
"Gender=Masc|NumType=Ord|Number=Plur|POS=ADJ",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Fut|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Number=Sing|POS=PRON",
"Number=Sing|POS=PRON|PronType=Dem",
"Mood=Ind|POS=VERB|VerbForm=Fin",
"Number=Plur|POS=DET|PronType=Dem",
"Gender=Masc|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Gender=Masc|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Gender=Masc|Number=Sing|POS=PRON",
"Gender=Masc|Number=Sing|POS=PRON|Person=3|PronType=Dem",
"Number=Sing|POS=PRON|Person=2|PronType=Prs",
"Gender=Masc|Number=Sing|POS=PRON|PronType=Rel",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Imp|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|NumType=Ord|Number=Sing|POS=ADJ",
"POS=PRON",
"POS=NUM",
"Gender=Fem|POS=NOUN",
"POS=SPACE",
"Gender=Fem|Number=Plur|POS=PRON",
"Number=Plur|POS=PRON|Person=3",
"Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Sing|POS=PRON|Person=1",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=PRON",
"Gender=Fem|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Mood=Sub|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"POS=INTJ",
"Number=Plur|POS=PRON|Person=2",
"NumType=Card|POS=PRON",
"Definite=Ind|Gender=Fem|Number=Plur|POS=DET|PronType=Art",
"Gender=Fem|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"NumType=Card|POS=NOUN",
"POS=PRON|PronType=Int",
"Gender=Fem|Number=Plur|POS=PRON|Person=3",
"Gender=Fem|Number=Sing|POS=DET",
"Mood=Cnd|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=DET",
"Mood=Sub|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Definite=Ind|Gender=Masc|Number=Plur|POS=DET|PronType=Art",
"Mood=Cnd|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=PRON|PronType=Dem",
"Gender=Masc|Number=Plur|POS=PROPN",
"Mood=Cnd|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=PRON|PronType=Dem",
"Number=Sing|POS=DET",
"Gender=Masc|NumType=Card|Number=Plur|POS=NOUN",
"Gender=Fem|Number=Plur|POS=PRON|PronType=Dem",
"Mood=Ind|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|POS=PRON",
"Gender=Masc|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Fem|Number=Sing|POS=PRON|PronType=Rel",
"Mood=Ind|Number=Sing|POS=AUX|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Cnd|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=AUX|Tense=Past|VerbForm=Part",
"POS=X",
"POS=SYM",
"Mood=Imp|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=DET|PronType=Int",
"Gender=Fem|Number=Plur|POS=DET|PronType=Int",
"POS=DET",
"Gender=Masc|Number=Plur|POS=PRON",
"Mood=Sub|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Mood=Ind|POS=VERB|Person=3|VerbForm=Fin",
"Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Mood=Cnd|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=2|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=DET|PronType=Int",
"Gender=Masc|Number=Plur|POS=DET",
"Gender=Fem|Number=Plur|POS=PRON|PronType=Rel",
"Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Masc|Number=Plur|POS=PRON|PronType=Rel",
"POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Fem|NumType=Ord|Number=Plur|POS=ADJ",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Fut|VerbForm=Fin",
"Mood=Imp|POS=VERB|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=2|Reflex=Yes",
"Mood=Cnd|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=1|Reflex=Yes",
"Gender=Masc|NumType=Card|Number=Sing|POS=NOUN",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Fut|VerbForm=Fin",
"Number=Sing|POS=PRON|Person=1|Reflex=Yes",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Gender=Masc|POS=PROPN",
"Mood=Cnd|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=1|PronType=Prs",
"Mood=Sub|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=2|PronType=Prs",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Fut|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Number=Sing|POS=PRON|Person=1|PronType=Prs",
"Mood=Cnd|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Plur|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Imp|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Plur|POS=AUX|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=2|Tense=Imp|VerbForm=Fin",
"Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Gender=Fem|Number=Plur|POS=PROPN",
"Gender=Masc|NumType=Card|POS=NUM"
],
"parser":[
"ROOT",
"acl",
"acl:relcl",
"advcl",
"advmod",
"amod",
"appos",
"aux:pass",
"aux:tense",
"case",
"cc",
"ccomp",
"conj",
"cop",
"dep",
"det",
"expl:comp",
"expl:pass",
"expl:subj",
"fixed",
"flat:foreign",
"flat:name",
"iobj",
"mark",
"nmod",
"nsubj",
"nsubj:pass",
"nummod",
"obj",
"obl:agent",
"obl:arg",
"obl:mod",
"parataxis",
"punct",
"vocative",
"xcomp"
],
"attribute_ruler":[
],
"lemmatizer":[
],
"ner":[
"LOC",
"MISC",
"ORG",
"PER"
]
},
"pipeline":[
"tok2vec",
"morphologizer",
"parser",
"attribute_ruler",
"lemmatizer",
"ner"
],
"components":[
"tok2vec",
"morphologizer",
"parser",
"senter",
"attribute_ruler",
"lemmatizer",
"ner"
],
"disabled":[
"senter"
],
"performance":{
"token_acc":0.997952498,
"token_p":0.9844389844,
"token_r":0.9896058454,
"token_f":0.9870156531,
"pos_acc":0.9617644028,
"morph_acc":0.9529502705,
"morph_micro_p":0.9796195652,
"morph_micro_r":0.9663701718,
"morph_micro_f":0.9729497638,
"morph_per_feat":{
"Definite":{
"p":0.986090776,
"r":0.9832116788,
"f":0.9846491228
},
"Number":{
"p":0.9897655378,
"r":0.979197349,
"f":0.9844530816
},
"PronType":{
"p":0.9935525467,
"r":0.9859245042,
"f":0.9897238279
},
"Gender":{
"p":0.9710519514,
"r":0.9601328904,
"f":0.9655615523
},
"Mood":{
"p":0.9597806216,
"r":0.9325044405,
"f":0.9459459459
},
"Person":{
"p":0.9818181818,
"r":0.9509433962,
"f":0.9661341853
},
"Tense":{
"p":0.9538784067,
"r":0.9295199183,
"f":0.9415416451
},
"VerbForm":{
"p":0.9686174724,
"r":0.9453642384,
"f":0.956849602
},
"NumType":{
"p":0.9790940767,
"r":0.9590443686,
"f":0.9689655172
},
"Reflex":{
"p":0.9772727273,
"r":0.9772727273,
"f":0.9772727273
},
"Voice":{
"p":0.9090909091,
"r":0.8928571429,
"f":0.9009009009
},
"Poss":{
"p":1.0,
"r":1.0,
"f":1.0
},
"Polarity":{
"p":1.0,
"r":0.9882352941,
"f":0.9940828402
}
},
"sents_p":0.8561151079,
"sents_r":0.8665048544,
"sents_f":0.861278649,
"dep_uas":0.8781984485,
"dep_las":0.8347514036,
"dep_las_per_type":{
"det":{
"p":0.9707078926,
"r":0.9628732849,
"f":0.9667747164
},
"nsubj":{
"p":0.853960396,
"r":0.8313253012,
"f":0.8424908425
},
"aux:tense":{
"p":0.905511811,
"r":0.92,
"f":0.9126984127
},
"root":{
"p":0.8490566038,
"r":0.8737864078,
"f":0.8612440191
},
"obj":{
"p":0.8179012346,
"r":0.7863501484,
"f":0.8018154312
},
"cc":{
"p":0.866359447,
"r":0.866359447,
"f":0.866359447
},
"case":{
"p":0.9595959596,
"r":0.9707084469,
"f":0.9651202167
},
"obl:mod":{
"p":0.6396103896,
"r":0.5880597015,
"f":0.6127527216
},
"nmod":{
"p":0.7996164909,
"r":0.8331668332,
"f":0.8160469667
},
"conj":{
"p":0.5,
"r":0.5,
"f":0.5
},
"nummod":{
"p":0.9079754601,
"r":0.875739645,
"f":0.8915662651
},
"amod":{
"p":0.8850364964,
"r":0.883424408,
"f":0.8842297174
},
"acl":{
"p":0.6445783133,
"r":0.6184971098,
"f":0.6312684366
},
"mark":{
"p":0.8590909091,
"r":0.8325991189,
"f":0.8456375839
},
"xcomp":{
"p":0.8068965517,
"r":0.7748344371,
"f":0.7905405405
},
"flat:name":{
"p":0.8691588785,
"r":0.8857142857,
"f":0.8773584906
},
"cop":{
"p":0.8636363636,
"r":0.8444444444,
"f":0.8539325843
},
"advmod":{
"p":0.8264984227,
"r":0.8213166144,
"f":0.8238993711
},
"obl:arg":{
"p":0.6409090909,
"r":0.6409090909,
"f":0.6409090909
},
"appos":{
"p":0.5063291139,
"r":0.4819277108,
"f":0.4938271605
},
"nsubj:pass":{
"p":0.8295454545,
"r":0.8588235294,
"f":0.8439306358
},
"aux:pass":{
"p":0.905982906,
"r":0.9464285714,
"f":0.9257641921
},
"acl:relcl":{
"p":0.5697674419,
"r":0.5697674419,
"f":0.5697674419
},
"advcl":{
"p":0.475,
"r":0.4871794872,
"f":0.4810126582
},
"fixed":{
"p":0.8372093023,
"r":0.72,
"f":0.7741935484
},
"dep":{
"p":0.2388059701,
"r":0.5517241379,
"f":0.3333333333
},
"expl:subj":{
"p":0.6857142857,
"r":0.75,
"f":0.7164179104
},
"expl:comp":{
"p":0.6388888889,
"r":0.7666666667,
"f":0.696969697
},
"expl:pass":{
"p":0.5,
"r":0.2857142857,
"f":0.3636363636
},
"ccomp":{
"p":0.7307692308,
"r":0.7450980392,
"f":0.7378640777
},
"obl:agent":{
"p":0.8717948718,
"r":0.8095238095,
"f":0.8395061728
},
"parataxis":{
"p":0.4285714286,
"r":0.3214285714,
"f":0.3673469388
},
"iobj":{
"p":0.75,
"r":0.6,
"f":0.6666666667
},
"nsubj:caus":{
"p":0.0,
"r":0.0,
"f":0.0
},
"aux:caus":{
"p":0.0,
"r":0.0,
"f":0.0
},
"obj:agent":{
"p":0.0,
"r":0.0,
"f":0.0
},
"goeswith":{
"p":0.0,
"r":0.0,
"f":0.0
},
"vocative":{
"p":0.8333333333,
"r":0.625,
"f":0.7142857143
},
"dislocated":{
"p":0.0,
"r":0.0,
"f":0.0
},
"flat:foreign":{
"p":0.0,
"r":0.0,
"f":0.0
},
"orphan":{
"p":0.0,
"r":0.0,
"f":0.0
},
"advcl:cleft":{
"p":0.0,
"r":0.0,
"f":0.0
},
"csubj":{
"p":0.0,
"r":0.0,
"f":0.0
}
},
"tag_acc":0.933216531,
"lemma_acc":0.9084463625,
"ents_p":0.8148438757,
"ents_r":0.8106360834,
"ents_f":0.8127345333,
"ents_per_type":{
"PER":{
"p":0.8671110481,
"r":0.8761195099,
"f":0.8715920026
},
"LOC":{
"p":0.8253083637,
"r":0.8424663264,
"f":0.8337990851
},
"ORG":{
"p":0.7624443545,
"r":0.7190839695,
"f":0.7401296405
},
"MISC":{
"p":0.6976186671,
"r":0.6363901443,
"f":0.6655992681
}
},
"speed":3569.3782435377
},
"sources":[
{
"name":"UD French Sequoia v2.8",
"url":"https://github.com/UniversalDependencies/UD_French-Sequoia",
"license":"LGPL-LR",
"author":"Candito, Marie; Seddah, Djam\u00e9; Perrier, Guy; Guillaume, Bruno"
},
{
"name":"WikiNER",
"url":"https://figshare.com/articles/Learning_multilingual_named_entity_recognition_from_Wikipedia/5462500",
"license":"CC BY 4.0",
"author":"Joel Nothman, Nicky Ringland, Will Radford, Tara Murphy, James R Curran"
},
{
"name":"spaCy lookups data",
"author":"Explosion",
"url":"https://github.com/explosion/spacy-lookups-data",
"license":"MIT"
}
],
"requirements":[
]
}