fr_core_news_sm / meta.json
Adriane Boyd
Update spaCy pipeline
6185031
raw history blame
No virus
18.9 kB
{
"lang":"fr",
"name":"core_news_sm",
"version":"3.4.0",
"description":"French pipeline optimized for CPU. Components: tok2vec, morphologizer, parser, senter, ner, attribute_ruler, lemmatizer.",
"author":"Explosion",
"email":"contact@explosion.ai",
"url":"https://explosion.ai",
"license":"LGPL-LR",
"spacy_version":">=3.4.0,<3.5.0",
"spacy_git_version":"dd038b536",
"vectors":{
"width":0,
"vectors":0,
"keys":0,
"name":null
},
"labels":{
"tok2vec":[
],
"morphologizer":[
"POS=PROPN",
"Gender=Fem|Number=Sing|POS=DET|PronType=Dem",
"Gender=Fem|Number=Sing|POS=NOUN",
"Number=Plur|POS=PRON|Person=1",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"POS=SCONJ",
"POS=ADP",
"Definite=Def|Gender=Masc|Number=Sing|POS=DET|PronType=Art",
"NumType=Ord|POS=ADJ",
"Gender=Masc|Number=Sing|POS=NOUN",
"POS=PUNCT",
"Gender=Masc|Number=Sing|POS=PROPN",
"Number=Plur|POS=ADJ",
"Gender=Masc|Number=Plur|POS=NOUN",
"Definite=Ind|Gender=Fem|Number=Sing|POS=DET|PronType=Art",
"Number=Sing|POS=ADJ",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Imp|VerbForm=Fin",
"POS=ADV",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Definite=Def|Gender=Fem|Number=Sing|POS=DET|PronType=Art",
"Gender=Fem|Number=Sing|POS=PROPN",
"Definite=Def|Number=Sing|POS=DET|PronType=Art",
"NumType=Card|POS=NUM",
"Definite=Def|Number=Plur|POS=DET|PronType=Art",
"Gender=Masc|Number=Plur|POS=ADJ",
"POS=CCONJ",
"Gender=Fem|Number=Plur|POS=NOUN",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Gender=Fem|Number=Plur|POS=ADJ",
"POS=ADJ",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"POS=PRON|PronType=Rel",
"Number=Sing|POS=DET|Poss=Yes",
"Definite=Def|Gender=Masc|Number=Sing|POS=ADP|PronType=Art",
"Definite=Def|Number=Plur|POS=ADP|PronType=Art",
"Definite=Ind|Number=Plur|POS=DET|PronType=Art",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Masc|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"POS=VERB|VerbForm=Inf",
"Gender=Fem|Number=Sing|POS=ADJ",
"Gender=Masc|Number=Sing|POS=PRON|Person=3",
"Number=Plur|POS=DET",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=ADJ",
"Gender=Masc|Number=Sing|POS=DET|PronType=Dem",
"POS=ADV|PronType=Int",
"POS=VERB|Tense=Pres|VerbForm=Part",
"Gender=Fem|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Ind|Gender=Masc|Number=Sing|POS=DET|PronType=Art",
"Gender=Masc|POS=ADJ",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Fut|VerbForm=Fin",
"Number=Plur|POS=DET|Poss=Yes",
"POS=AUX|VerbForm=Inf",
"Gender=Masc|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Masc|POS=VERB|Tense=Past|VerbForm=Part",
"POS=ADV|Polarity=Neg",
"Definite=Ind|Number=Sing|POS=DET|PronType=Art",
"Gender=Fem|Number=Sing|POS=PRON|Person=3",
"POS=PRON|Person=3|Reflex=Yes",
"Gender=Masc|POS=NOUN",
"POS=AUX|Tense=Past|VerbForm=Part",
"POS=PRON|Person=3",
"Number=Plur|POS=NOUN",
"NumType=Ord|Number=Sing|POS=ADJ",
"POS=VERB|Tense=Past|VerbForm=Part",
"POS=AUX|Tense=Pres|VerbForm=Part",
"Gender=Masc|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Sing|POS=PRON|Person=3",
"Number=Sing|POS=NOUN",
"Gender=Masc|Number=Plur|POS=PRON|Person=3",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Imp|VerbForm=Fin",
"Gender=Fem|NumType=Ord|Number=Sing|POS=ADJ",
"Number=Plur|POS=PROPN",
"Number=Sing|POS=PROPN",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Plur|POS=PRON|PronType=Dem",
"Gender=Masc|Number=Sing|POS=DET",
"Gender=Fem|Number=Sing|POS=DET|Poss=Yes",
"Gender=Masc|POS=PRON",
"POS=NOUN",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON",
"Gender=Masc|NumType=Ord|Number=Plur|POS=ADJ",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Fut|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Number=Sing|POS=PRON",
"Number=Sing|POS=PRON|PronType=Dem",
"Mood=Ind|POS=VERB|VerbForm=Fin",
"Number=Plur|POS=DET|PronType=Dem",
"Gender=Masc|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Gender=Masc|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Gender=Masc|Number=Sing|POS=PRON",
"Gender=Masc|Number=Sing|POS=PRON|Person=3|PronType=Dem",
"Number=Sing|POS=PRON|Person=2|PronType=Prs",
"Gender=Masc|Number=Sing|POS=PRON|PronType=Rel",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Imp|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|NumType=Ord|Number=Sing|POS=ADJ",
"POS=PRON",
"POS=NUM",
"Gender=Fem|POS=NOUN",
"POS=SPACE",
"Gender=Fem|Number=Plur|POS=PRON",
"Number=Plur|POS=PRON|Person=3",
"Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Sing|POS=PRON|Person=1",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=PRON",
"Gender=Fem|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Mood=Sub|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"POS=INTJ",
"Number=Plur|POS=PRON|Person=2",
"NumType=Card|POS=PRON",
"Definite=Ind|Gender=Fem|Number=Plur|POS=DET|PronType=Art",
"Gender=Fem|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"NumType=Card|POS=NOUN",
"POS=PRON|PronType=Int",
"Gender=Fem|Number=Plur|POS=PRON|Person=3",
"Gender=Fem|Number=Sing|POS=DET",
"Mood=Cnd|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=DET",
"Mood=Sub|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Definite=Ind|Gender=Masc|Number=Plur|POS=DET|PronType=Art",
"Mood=Cnd|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=PRON|PronType=Dem",
"Gender=Masc|Number=Plur|POS=PROPN",
"Mood=Cnd|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=PRON|PronType=Dem",
"Number=Sing|POS=DET",
"Gender=Masc|NumType=Card|Number=Plur|POS=NOUN",
"Gender=Fem|Number=Plur|POS=PRON|PronType=Dem",
"Mood=Ind|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|POS=PRON",
"Gender=Masc|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Fem|Number=Sing|POS=PRON|PronType=Rel",
"Mood=Ind|Number=Sing|POS=AUX|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Cnd|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=AUX|Tense=Past|VerbForm=Part",
"POS=X",
"POS=SYM",
"Mood=Imp|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=DET|PronType=Int",
"Gender=Fem|Number=Plur|POS=DET|PronType=Int",
"POS=DET",
"Gender=Masc|Number=Plur|POS=PRON",
"Mood=Sub|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Mood=Ind|POS=VERB|Person=3|VerbForm=Fin",
"Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Mood=Cnd|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=2|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=DET|PronType=Int",
"Gender=Masc|Number=Plur|POS=DET",
"Gender=Fem|Number=Plur|POS=PRON|PronType=Rel",
"Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Masc|Number=Plur|POS=PRON|PronType=Rel",
"POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Fem|NumType=Ord|Number=Plur|POS=ADJ",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Fut|VerbForm=Fin",
"Mood=Imp|POS=VERB|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=2|Reflex=Yes",
"Mood=Cnd|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=1|Reflex=Yes",
"Gender=Masc|NumType=Card|Number=Sing|POS=NOUN",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Fut|VerbForm=Fin",
"Number=Sing|POS=PRON|Person=1|Reflex=Yes",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Gender=Masc|POS=PROPN",
"Mood=Cnd|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=1|PronType=Prs",
"Mood=Sub|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=2|PronType=Prs",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Fut|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Number=Sing|POS=PRON|Person=1|PronType=Prs",
"Mood=Cnd|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Plur|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Imp|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Plur|POS=AUX|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=2|Tense=Imp|VerbForm=Fin",
"Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Gender=Fem|Number=Plur|POS=PROPN",
"Gender=Masc|NumType=Card|POS=NUM"
],
"parser":[
"ROOT",
"acl",
"acl:relcl",
"advcl",
"advmod",
"amod",
"appos",
"aux:pass",
"aux:tense",
"case",
"cc",
"ccomp",
"conj",
"cop",
"dep",
"det",
"expl:comp",
"expl:pass",
"expl:subj",
"fixed",
"flat:foreign",
"flat:name",
"iobj",
"mark",
"nmod",
"nsubj",
"nsubj:pass",
"nummod",
"obj",
"obl:agent",
"obl:arg",
"obl:mod",
"parataxis",
"punct",
"vocative",
"xcomp"
],
"attribute_ruler":[
],
"lemmatizer":[
],
"ner":[
"LOC",
"MISC",
"ORG",
"PER"
]
},
"pipeline":[
"tok2vec",
"morphologizer",
"parser",
"attribute_ruler",
"lemmatizer",
"ner"
],
"components":[
"tok2vec",
"morphologizer",
"parser",
"senter",
"attribute_ruler",
"lemmatizer",
"ner"
],
"disabled":[
"senter"
],
"performance":{
"token_acc":0.9989751998,
"token_p":0.9844389844,
"token_r":0.9896058454,
"token_f":0.9870156531,
"pos_acc":0.9620735855,
"morph_acc":0.952582208,
"morph_micro_p":0.978246133,
"morph_micro_r":0.967101255,
"morph_micro_f":0.9726417696,
"morph_per_feat":{
"Definite":{
"p":0.9846378932,
"r":0.9824817518,
"f":0.9835586408
},
"Number":{
"p":0.9907063197,
"r":0.9812223859,
"f":0.9859415464
},
"PronType":{
"p":0.995483871,
"r":0.9872040947,
"f":0.9913266945
},
"Gender":{
"p":0.9706563707,
"r":0.9637107079,
"f":0.9671710695
},
"Mood":{
"p":0.9539594843,
"r":0.920071048,
"f":0.9367088608
},
"Person":{
"p":0.9779507134,
"r":0.948427673,
"f":0.962962963
},
"Tense":{
"p":0.9447916667,
"r":0.9264555669,
"f":0.9355337803
},
"VerbForm":{
"p":0.9586846543,
"r":0.9412251656,
"f":0.9498746867
},
"NumType":{
"p":0.9929328622,
"r":0.9590443686,
"f":0.9756944444
},
"Reflex":{
"p":1.0,
"r":1.0,
"f":1.0
},
"Voice":{
"p":0.8793103448,
"r":0.9107142857,
"f":0.8947368421
},
"Poss":{
"p":1.0,
"r":1.0,
"f":1.0
},
"Polarity":{
"p":0.9882352941,
"r":0.9882352941,
"f":0.9882352941
}
},
"sents_p":0.8677884615,
"sents_r":0.8791762895,
"sents_f":0.8719806763,
"dep_uas":0.8742764529,
"dep_las":0.8295618959,
"dep_las_per_type":{
"det":{
"p":0.9675587997,
"r":0.9628732849,
"f":0.965210356
},
"nsubj":{
"p":0.8341463415,
"r":0.8240963855,
"f":0.8290909091
},
"aux:tense":{
"p":0.8976377953,
"r":0.912,
"f":0.9047619048
},
"root":{
"p":0.8474576271,
"r":0.8495145631,
"f":0.8484848485
},
"obj":{
"p":0.8058823529,
"r":0.8130563798,
"f":0.8094534712
},
"cc":{
"p":0.8630136986,
"r":0.8709677419,
"f":0.8669724771
},
"case":{
"p":0.9569603228,
"r":0.969346049,
"f":0.9631133672
},
"obl:mod":{
"p":0.6314102564,
"r":0.5880597015,
"f":0.6089644513
},
"nmod":{
"p":0.7830985915,
"r":0.8331668332,
"f":0.807357212
},
"conj":{
"p":0.4836065574,
"r":0.4645669291,
"f":0.4738955823
},
"nummod":{
"p":0.917721519,
"r":0.8579881657,
"f":0.8868501529
},
"amod":{
"p":0.8621323529,
"r":0.85428051,
"f":0.8581884721
},
"acl":{
"p":0.6503067485,
"r":0.612716763,
"f":0.630952381
},
"mark":{
"p":0.8812785388,
"r":0.8502202643,
"f":0.865470852
},
"xcomp":{
"p":0.7388535032,
"r":0.7682119205,
"f":0.7532467532
},
"flat:name":{
"p":0.8878504673,
"r":0.9047619048,
"f":0.8962264151
},
"cop":{
"p":0.8505747126,
"r":0.8222222222,
"f":0.8361581921
},
"advmod":{
"p":0.8161290323,
"r":0.7931034483,
"f":0.8044515103
},
"obl:arg":{
"p":0.6601941748,
"r":0.6181818182,
"f":0.6384976526
},
"appos":{
"p":0.4444444444,
"r":0.4337349398,
"f":0.4390243902
},
"nsubj:pass":{
"p":0.8292682927,
"r":0.8,
"f":0.8143712575
},
"aux:pass":{
"p":0.9035087719,
"r":0.9196428571,
"f":0.9115044248
},
"acl:relcl":{
"p":0.6626506024,
"r":0.6395348837,
"f":0.650887574
},
"advcl":{
"p":0.5057471264,
"r":0.5641025641,
"f":0.5333333333
},
"fixed":{
"p":0.8452380952,
"r":0.71,
"f":0.7717391304
},
"dep":{
"p":0.2777777778,
"r":0.5172413793,
"f":0.3614457831
},
"expl:subj":{
"p":0.7647058824,
"r":0.8125,
"f":0.7878787879
},
"expl:comp":{
"p":0.6666666667,
"r":0.8666666667,
"f":0.7536231884
},
"expl:pass":{
"p":0.25,
"r":0.1428571429,
"f":0.1818181818
},
"obl:agent":{
"p":0.8139534884,
"r":0.8333333333,
"f":0.8235294118
},
"ccomp":{
"p":0.7391304348,
"r":0.6666666667,
"f":0.7010309278
},
"parataxis":{
"p":0.5416666667,
"r":0.4642857143,
"f":0.5
},
"iobj":{
"p":0.6,
"r":0.48,
"f":0.5333333333
},
"nsubj:caus":{
"p":0.0,
"r":0.0,
"f":0.0
},
"aux:caus":{
"p":0.0,
"r":0.0,
"f":0.0
},
"obj:agent":{
"p":0.0,
"r":0.0,
"f":0.0
},
"goeswith":{
"p":0.0,
"r":0.0,
"f":0.0
},
"vocative":{
"p":1.0,
"r":0.625,
"f":0.7692307692
},
"dislocated":{
"p":0.0,
"r":0.0,
"f":0.0
},
"flat:foreign":{
"p":1.0,
"r":0.2857142857,
"f":0.4444444444
},
"orphan":{
"p":0.0,
"r":0.0,
"f":0.0
},
"advcl:cleft":{
"p":0.0,
"r":0.0,
"f":0.0
},
"csubj":{
"p":0.0,
"r":0.0,
"f":0.0
}
},
"tag_acc":0.9334226528,
"lemma_acc":0.9028976572,
"ents_p":0.812993431,
"ents_r":0.8100156119,
"ents_f":0.8115017896,
"ents_per_type":{
"LOC":{
"p":0.8247443931,
"r":0.8368192086,
"f":0.8307379262
},
"PER":{
"p":0.8687411598,
"r":0.8801318335,
"f":0.8743994021
},
"ORG":{
"p":0.7622446956,
"r":0.7335877863,
"f":0.7476417388
},
"MISC":{
"p":0.6840694006,
"r":0.6323079166,
"f":0.6571709978
}
},
"speed":4128.958361591
},
"sources":[
{
"name":"UD French Sequoia v2.8",
"url":"https://github.com/UniversalDependencies/UD_French-Sequoia",
"license":"LGPL-LR",
"author":"Candito, Marie; Seddah, Djam\u00e9; Perrier, Guy; Guillaume, Bruno"
},
{
"name":"WikiNER",
"url":"https://figshare.com/articles/Learning_multilingual_named_entity_recognition_from_Wikipedia/5462500",
"license":"CC BY 4.0",
"author":"Joel Nothman, Nicky Ringland, Will Radford, Tara Murphy, James R Curran"
},
{
"name":"spaCy lookups data",
"author":"Explosion",
"url":"https://github.com/explosion/spacy-lookups-data",
"license":"MIT"
}
],
"requirements":[
]
}