fr_core_news_md / meta.json
EC2 Default User
Update spaCy pipeline
c636104
raw history blame
No virus
19.1 kB
{
"lang":"fr",
"name":"core_news_md",
"version":"3.3.0",
"description":"French pipeline optimized for CPU. Components: tok2vec, morphologizer, parser, senter, ner, attribute_ruler, lemmatizer.",
"author":"Explosion",
"email":"contact@explosion.ai",
"url":"https://explosion.ai",
"license":"LGPL-LR",
"spacy_version":">=3.3.0.dev0,<3.4.0",
"spacy_git_version":"849bef2de",
"vectors":{
"width":300,
"vectors":20000,
"keys":500000,
"name":"fr_vectors"
},
"labels":{
"tok2vec":[
],
"morphologizer":[
"POS=PROPN",
"Gender=Fem|Number=Sing|POS=DET|PronType=Dem",
"Gender=Fem|Number=Sing|POS=NOUN",
"Number=Plur|POS=PRON|Person=1",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"POS=SCONJ",
"POS=ADP",
"Definite=Def|Gender=Masc|Number=Sing|POS=DET|PronType=Art",
"NumType=Ord|POS=ADJ",
"Gender=Masc|Number=Sing|POS=NOUN",
"POS=PUNCT",
"Gender=Masc|Number=Sing|POS=PROPN",
"Number=Plur|POS=ADJ",
"Gender=Masc|Number=Plur|POS=NOUN",
"Definite=Ind|Gender=Fem|Number=Sing|POS=DET|PronType=Art",
"Number=Sing|POS=ADJ",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Imp|VerbForm=Fin",
"POS=ADV",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Definite=Def|Gender=Fem|Number=Sing|POS=DET|PronType=Art",
"Gender=Fem|Number=Sing|POS=PROPN",
"Definite=Def|Number=Sing|POS=DET|PronType=Art",
"NumType=Card|POS=NUM",
"Definite=Def|Number=Plur|POS=DET|PronType=Art",
"Gender=Masc|Number=Plur|POS=ADJ",
"POS=CCONJ",
"Gender=Fem|Number=Plur|POS=NOUN",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Gender=Fem|Number=Plur|POS=ADJ",
"POS=ADJ",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"POS=PRON|PronType=Rel",
"Number=Sing|POS=DET|Poss=Yes",
"Definite=Def|Gender=Masc|Number=Sing|POS=ADP|PronType=Art",
"Definite=Def|Number=Plur|POS=ADP|PronType=Art",
"Definite=Ind|Number=Plur|POS=DET|PronType=Art",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Masc|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"POS=VERB|VerbForm=Inf",
"Gender=Fem|Number=Sing|POS=ADJ",
"Gender=Masc|Number=Sing|POS=PRON|Person=3",
"Number=Plur|POS=DET",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=ADJ",
"Gender=Masc|Number=Sing|POS=DET|PronType=Dem",
"POS=ADV|PronType=Int",
"POS=VERB|Tense=Pres|VerbForm=Part",
"Gender=Fem|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Ind|Gender=Masc|Number=Sing|POS=DET|PronType=Art",
"Gender=Masc|POS=ADJ",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Fut|VerbForm=Fin",
"Number=Plur|POS=DET|Poss=Yes",
"POS=AUX|VerbForm=Inf",
"Gender=Masc|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Masc|POS=VERB|Tense=Past|VerbForm=Part",
"POS=ADV|Polarity=Neg",
"Definite=Ind|Number=Sing|POS=DET|PronType=Art",
"Gender=Fem|Number=Sing|POS=PRON|Person=3",
"POS=PRON|Person=3|Reflex=Yes",
"Gender=Masc|POS=NOUN",
"POS=AUX|Tense=Past|VerbForm=Part",
"POS=PRON|Person=3",
"Number=Plur|POS=NOUN",
"NumType=Ord|Number=Sing|POS=ADJ",
"POS=VERB|Tense=Past|VerbForm=Part",
"POS=AUX|Tense=Pres|VerbForm=Part",
"Gender=Masc|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Sing|POS=PRON|Person=3",
"Number=Sing|POS=NOUN",
"Gender=Masc|Number=Plur|POS=PRON|Person=3",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Imp|VerbForm=Fin",
"Gender=Fem|NumType=Ord|Number=Sing|POS=ADJ",
"Number=Plur|POS=PROPN",
"Number=Sing|POS=PROPN",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Plur|POS=PRON|PronType=Dem",
"Gender=Masc|Number=Sing|POS=DET",
"Gender=Fem|Number=Sing|POS=DET|Poss=Yes",
"Gender=Masc|POS=PRON",
"POS=NOUN",
"Mood=Ind|Number=Sing|POS=VERB|Person=3|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=3|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON",
"Gender=Masc|NumType=Ord|Number=Plur|POS=ADJ",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Fut|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Number=Sing|POS=PRON",
"Number=Sing|POS=PRON|PronType=Dem",
"Mood=Ind|POS=VERB|VerbForm=Fin",
"Number=Plur|POS=DET|PronType=Dem",
"Gender=Masc|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Gender=Masc|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Gender=Masc|Number=Sing|POS=PRON",
"Gender=Masc|Number=Sing|POS=PRON|Person=3|PronType=Dem",
"Number=Sing|POS=PRON|Person=2|PronType=Prs",
"Gender=Masc|Number=Sing|POS=PRON|PronType=Rel",
"Mood=Ind|Number=Plur|POS=AUX|Person=3|Tense=Imp|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|NumType=Ord|Number=Sing|POS=ADJ",
"POS=PRON",
"POS=NUM",
"Gender=Fem|POS=NOUN",
"Gender=Fem|Number=Plur|POS=PRON",
"Number=Plur|POS=PRON|Person=3",
"Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Sing|POS=PRON|Person=1",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=VERB|Person=3|Tense=Past|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=PRON",
"Gender=Fem|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Mood=Sub|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"POS=INTJ",
"Number=Plur|POS=PRON|Person=2",
"NumType=Card|POS=PRON",
"Definite=Ind|Gender=Fem|Number=Plur|POS=DET|PronType=Art",
"Gender=Fem|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"NumType=Card|POS=NOUN",
"POS=PRON|PronType=Int",
"Gender=Fem|Number=Plur|POS=PRON|Person=3",
"Gender=Fem|Number=Sing|POS=DET",
"Mood=Cnd|Number=Sing|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=DET",
"Mood=Sub|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Definite=Ind|Gender=Masc|Number=Plur|POS=DET|PronType=Art",
"Mood=Cnd|Number=Sing|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=PRON|PronType=Dem",
"Gender=Masc|Number=Plur|POS=PROPN",
"Mood=Cnd|Number=Plur|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=PRON|PronType=Dem",
"Number=Sing|POS=DET",
"Gender=Masc|NumType=Card|Number=Plur|POS=NOUN",
"Gender=Fem|Number=Plur|POS=PRON|PronType=Dem",
"Mood=Ind|POS=VERB|Person=3|Tense=Pres|VerbForm=Fin",
"Gender=Fem|POS=PRON",
"Gender=Masc|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Fem|Number=Sing|POS=PRON|PronType=Rel",
"Mood=Ind|Number=Sing|POS=AUX|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Cnd|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=AUX|Tense=Past|VerbForm=Part",
"POS=X",
"POS=SYM",
"Mood=Imp|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Gender=Masc|Number=Sing|POS=DET|PronType=Int",
"Gender=Fem|Number=Plur|POS=DET|PronType=Int",
"POS=DET",
"Gender=Masc|Number=Plur|POS=PRON",
"Mood=Sub|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Mood=Ind|POS=VERB|Person=3|VerbForm=Fin",
"Number=Sing|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Mood=Cnd|Number=Plur|POS=VERB|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=2|Tense=Pres|VerbForm=Fin",
"Gender=Fem|Number=Sing|POS=DET|PronType=Int",
"Gender=Masc|Number=Plur|POS=DET",
"Gender=Fem|Number=Plur|POS=PRON|PronType=Rel",
"Number=Plur|POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Masc|Number=Plur|POS=PRON|PronType=Rel",
"POS=VERB|Tense=Past|VerbForm=Part|Voice=Pass",
"Gender=Fem|NumType=Ord|Number=Plur|POS=ADJ",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Fut|VerbForm=Fin",
"Mood=Imp|POS=VERB|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=2|Reflex=Yes",
"Mood=Cnd|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=1|Reflex=Yes",
"Gender=Masc|NumType=Card|Number=Sing|POS=NOUN",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Fut|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Fut|VerbForm=Fin",
"Number=Sing|POS=PRON|Person=1|Reflex=Yes",
"Mood=Ind|Number=Plur|POS=VERB|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=AUX|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Imp|VerbForm=Fin",
"Mood=Sub|Number=Sing|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Gender=Masc|POS=PROPN",
"Mood=Cnd|Number=Plur|POS=AUX|Person=3|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=1|PronType=Prs",
"Mood=Sub|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Number=Plur|POS=PRON|Person=2|PronType=Prs",
"Mood=Ind|Number=Sing|POS=VERB|Person=1|Tense=Fut|VerbForm=Fin",
"Gender=Fem|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Number=Sing|POS=PRON|Person=1|PronType=Prs",
"Mood=Cnd|Number=Sing|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Plur|POS=AUX|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Imp|Number=Plur|POS=VERB|Person=1|Tense=Pres|VerbForm=Fin",
"Mood=Sub|Number=Plur|POS=AUX|Person=2|Tense=Pres|VerbForm=Fin",
"Mood=Ind|Number=Plur|POS=VERB|Person=2|Tense=Imp|VerbForm=Fin",
"Mood=Ind|Number=Sing|POS=AUX|Person=2|Tense=Imp|VerbForm=Fin",
"Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Gender=Fem|Number=Plur|POS=PROPN",
"Gender=Masc|NumType=Card|POS=NUM"
],
"parser":[
"ROOT",
"acl",
"acl:relcl",
"advcl",
"advmod",
"amod",
"appos",
"aux:pass",
"aux:tense",
"case",
"cc",
"ccomp",
"conj",
"cop",
"dep",
"det",
"expl:comp",
"expl:pass",
"expl:subj",
"fixed",
"flat:foreign",
"flat:name",
"iobj",
"mark",
"nmod",
"nsubj",
"nsubj:pass",
"nummod",
"obj",
"obl:agent",
"obl:arg",
"obl:mod",
"parataxis",
"punct",
"vocative",
"xcomp"
],
"attribute_ruler":[
],
"lemmatizer":[
],
"ner":[
"LOC",
"MISC",
"ORG",
"PER"
]
},
"pipeline":[
"tok2vec",
"morphologizer",
"parser",
"attribute_ruler",
"lemmatizer",
"ner"
],
"components":[
"tok2vec",
"morphologizer",
"parser",
"senter",
"attribute_ruler",
"lemmatizer",
"ner"
],
"disabled":[
"senter"
],
"performance":{
"token_acc":0.9989751998,
"token_p":0.9844389844,
"token_r":0.9896058454,
"token_f":0.9870156531,
"pos_acc":0.9710899253,
"morph_acc":0.9620122674,
"morph_micro_p":0.9839828168,
"morph_micro_r":0.9768490313,
"morph_micro_f":0.9804029472,
"morph_per_feat":{
"Definite":{
"p":0.9868613139,
"r":0.9868613139,
"f":0.9868613139
},
"Number":{
"p":0.992584353,
"r":0.985640648,
"f":0.9891003141
},
"PronType":{
"p":0.9942159383,
"r":0.9897632758,
"f":0.9919846105
},
"Gender":{
"p":0.9798813516,
"r":0.970866343,
"f":0.9753530167
},
"Mood":{
"p":0.9675090253,
"r":0.9520426288,
"f":0.9597135184
},
"Person":{
"p":0.9821428571,
"r":0.9685534591,
"f":0.9753008233
},
"Tense":{
"p":0.9578622816,
"r":0.9519918284,
"f":0.9549180328
},
"VerbForm":{
"p":0.9742310889,
"r":0.9701986755,
"f":0.972210701
},
"NumType":{
"p":0.9893992933,
"r":0.9556313993,
"f":0.9722222222
},
"Reflex":{
"p":1.0,
"r":1.0,
"f":1.0
},
"Voice":{
"p":0.9304347826,
"r":0.9553571429,
"f":0.9427312775
},
"Poss":{
"p":0.9827586207,
"r":1.0,
"f":0.9913043478
},
"Polarity":{
"p":1.0,
"r":0.9882352941,
"f":0.9940828402
}
},
"sents_p":0.853427896,
"sents_r":0.8919553503,
"sents_f":0.8646706587,
"dep_uas":0.8977489729,
"dep_las":0.8591264102,
"dep_las_per_type":{
"det":{
"p":0.9854014599,
"r":0.98062954,
"f":0.9830097087
},
"nsubj":{
"p":0.8819277108,
"r":0.8819277108,
"f":0.8819277108
},
"aux:tense":{
"p":0.952,
"r":0.952,
"f":0.952
},
"root":{
"p":0.8717339667,
"r":0.890776699,
"f":0.881152461
},
"obj":{
"p":0.8700906344,
"r":0.8545994065,
"f":0.8622754491
},
"cc":{
"p":0.8986175115,
"r":0.8986175115,
"f":0.8986175115
},
"case":{
"p":0.9695328368,
"r":0.9754768392,
"f":0.9724957555
},
"obl:mod":{
"p":0.6677316294,
"r":0.623880597,
"f":0.6450617284
},
"nmod":{
"p":0.8051330798,
"r":0.8461538462,
"f":0.8251339503
},
"conj":{
"p":0.5863453815,
"r":0.5748031496,
"f":0.5805168986
},
"nummod":{
"p":0.9141104294,
"r":0.8816568047,
"f":0.8975903614
},
"amod":{
"p":0.9037037037,
"r":0.8888888889,
"f":0.8962350781
},
"acl":{
"p":0.6892655367,
"r":0.7052023121,
"f":0.6971428571
},
"mark":{
"p":0.88,
"r":0.872246696,
"f":0.8761061947
},
"xcomp":{
"p":0.8476821192,
"r":0.8476821192,
"f":0.8476821192
},
"flat:name":{
"p":0.9223300971,
"r":0.9047619048,
"f":0.9134615385
},
"cop":{
"p":0.8709677419,
"r":0.9,
"f":0.8852459016
},
"advmod":{
"p":0.8698412698,
"r":0.8589341693,
"f":0.8643533123
},
"obl:arg":{
"p":0.6872037915,
"r":0.6590909091,
"f":0.6728538283
},
"appos":{
"p":0.5263157895,
"r":0.4819277108,
"f":0.5031446541
},
"nsubj:pass":{
"p":0.9156626506,
"r":0.8941176471,
"f":0.9047619048
},
"aux:pass":{
"p":0.963963964,
"r":0.9553571429,
"f":0.9596412556
},
"acl:relcl":{
"p":0.6511627907,
"r":0.6511627907,
"f":0.6511627907
},
"advcl":{
"p":0.5128205128,
"r":0.5128205128,
"f":0.5128205128
},
"fixed":{
"p":0.8539325843,
"r":0.76,
"f":0.8042328042
},
"dep":{
"p":0.2571428571,
"r":0.6206896552,
"f":0.3636363636
},
"expl:subj":{
"p":0.8666666667,
"r":0.8125,
"f":0.8387096774
},
"expl:comp":{
"p":0.625,
"r":0.8333333333,
"f":0.7142857143
},
"expl:pass":{
"p":0.25,
"r":0.1428571429,
"f":0.1818181818
},
"ccomp":{
"p":0.6666666667,
"r":0.6666666667,
"f":0.6666666667
},
"parataxis":{
"p":0.5714285714,
"r":0.4285714286,
"f":0.4897959184
},
"iobj":{
"p":0.75,
"r":0.48,
"f":0.5853658537
},
"obl:agent":{
"p":0.8918918919,
"r":0.7857142857,
"f":0.835443038
},
"nsubj:caus":{
"p":0.0,
"r":0.0,
"f":0.0
},
"aux:caus":{
"p":0.0,
"r":0.0,
"f":0.0
},
"obj:agent":{
"p":0.0,
"r":0.0,
"f":0.0
},
"goeswith":{
"p":0.0,
"r":0.0,
"f":0.0
},
"vocative":{
"p":1.0,
"r":0.625,
"f":0.7692307692
},
"dislocated":{
"p":0.0,
"r":0.0,
"f":0.0
},
"flat:foreign":{
"p":0.0,
"r":0.0,
"f":0.0
},
"orphan":{
"p":0.0,
"r":0.0,
"f":0.0
},
"advcl:cleft":{
"p":0.0,
"r":0.0,
"f":0.0
},
"csubj":{
"p":0.0,
"r":0.0,
"f":0.0
}
},
"tag_acc":0.9424375161,
"lemma_acc":0.9063912865,
"ents_p":0.8317787005,
"ents_r":0.8299307474,
"ents_f":0.8308536964,
"ents_per_type":{
"PER":{
"p":0.8974413195,
"r":0.912230422,
"f":0.9047754406
},
"LOC":{
"p":0.8422664625,
"r":0.851250732,
"f":0.8467347661
},
"ORG":{
"p":0.7548623147,
"r":0.7480916031,
"f":0.751461708
},
"MISC":{
"p":0.7072900158,
"r":0.6506779414,
"f":0.6778039335
}
},
"speed":4445.2953793177
},
"sources":[
{
"name":"UD French Sequoia v2.8",
"url":"https://github.com/UniversalDependencies/UD_French-Sequoia",
"license":"LGPL-LR",
"author":"Candito, Marie; Seddah, Djam\u00e9; Perrier, Guy; Guillaume, Bruno"
},
{
"name":"WikiNER",
"url":"https://figshare.com/articles/Learning_multilingual_named_entity_recognition_from_Wikipedia/5462500",
"license":"CC BY 4.0",
"author":"Joel Nothman, Nicky Ringland, Will Radford, Tara Murphy, James R Curran"
},
{
"name":"spaCy lookups data",
"author":"Explosion",
"url":"https://github.com/explosion/spacy-lookups-data",
"license":"MIT"
},
{
"name":"Explosion fastText Vectors (cbow, OSCAR Common Crawl + Wikipedia)",
"url":"https://spacy.io",
"license":"CC0",
"author":"Explosion"
}
],
"requirements":[
]
}