da_core_news_lg / meta.json
adrianeboyd's picture
Update spaCy pipeline
a2f0070
{
"lang":"da",
"name":"core_news_lg",
"version":"3.7.0",
"description":"Danish pipeline optimized for CPU. Components: tok2vec, morphologizer, parser, lemmatizer (trainable_lemmatizer), senter, ner, attribute_ruler.",
"author":"Explosion",
"email":"contact@explosion.ai",
"url":"https://explosion.ai",
"license":"CC BY-SA 4.0",
"spacy_version":">=3.7.0,<3.8.0",
"spacy_git_version":"6b4f77441",
"vectors":{
"width":300,
"vectors":500000,
"keys":500000,
"name":"da_vectors"
},
"labels":{
"tok2vec":[
],
"morphologizer":[
"AdpType=Prep|POS=ADP",
"Definite=Ind|Gender=Com|Number=Sing|POS=NOUN",
"Mood=Ind|POS=AUX|Tense=Pres|VerbForm=Fin|Voice=Act",
"POS=PROPN",
"Definite=Ind|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Def|Gender=Neut|Number=Sing|POS=NOUN",
"POS=SCONJ",
"Definite=Def|Gender=Com|Number=Sing|POS=NOUN",
"Mood=Ind|POS=VERB|Tense=Pres|VerbForm=Fin|Voice=Act",
"POS=ADV",
"Number=Plur|POS=DET|PronType=Dem",
"Degree=Pos|Number=Plur|POS=ADJ",
"Definite=Ind|Gender=Com|Number=Plur|POS=NOUN",
"POS=PUNCT",
"POS=CCONJ",
"Definite=Ind|Degree=Cmp|Number=Sing|POS=ADJ",
"Degree=Cmp|POS=ADJ",
"POS=PRON|PartType=Inf",
"Gender=Com|Number=Sing|POS=DET|PronType=Ind",
"Definite=Ind|Degree=Pos|Number=Sing|POS=ADJ",
"Case=Acc|Gender=Neut|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Definite=Ind|Gender=Neut|Number=Plur|POS=NOUN",
"Definite=Def|Degree=Pos|Number=Sing|POS=ADJ",
"Gender=Neut|Number=Sing|POS=DET|PronType=Dem",
"Degree=Pos|POS=ADV",
"Definite=Def|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Ind|Gender=Neut|Number=Sing|POS=NOUN",
"POS=PRON|PronType=Dem",
"NumType=Card|POS=NUM",
"Definite=Ind|Degree=Pos|Gender=Neut|Number=Sing|POS=ADJ",
"Case=Acc|Gender=Com|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Degree=Pos|Gender=Com|Number=Sing|POS=ADJ",
"Case=Nom|Gender=Com|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"NumType=Ord|POS=ADJ",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=DET|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Mood=Ind|POS=AUX|Tense=Past|VerbForm=Fin|Voice=Act",
"POS=VERB|VerbForm=Inf|Voice=Act",
"Mood=Ind|POS=VERB|Tense=Past|VerbForm=Fin|Voice=Act",
"POS=NOUN",
"Mood=Ind|POS=VERB|Tense=Pres|VerbForm=Fin|Voice=Pass",
"POS=ADP|PartType=Inf",
"Degree=Pos|POS=ADJ",
"Definite=Def|Gender=Com|Number=Plur|POS=NOUN",
"Number[psor]=Sing|POS=DET|Person=3|Poss=Yes|PronType=Prs",
"Case=Gen|Definite=Def|Gender=Com|Number=Sing|POS=NOUN",
"POS=AUX|VerbForm=Inf|Voice=Act",
"Definite=Ind|Degree=Pos|Gender=Com|Number=Sing|POS=ADJ",
"Gender=Com|Number=Sing|POS=DET|PronType=Dem",
"Number=Plur|POS=DET|PronType=Ind",
"Gender=Com|Number=Sing|POS=PRON|PronType=Ind",
"Case=Acc|POS=PRON|Person=3|PronType=Prs|Reflex=Yes",
"POS=PART|PartType=Inf",
"Gender=Neut|Number=Sing|POS=DET|PronType=Ind",
"Case=Acc|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Case=Gen|Definite=Def|Gender=Neut|Number=Sing|POS=NOUN",
"Case=Nom|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Case=Nom|Gender=Com|Number=Sing|POS=PRON|Person=1|PronType=Prs",
"Case=Nom|Gender=Com|POS=PRON|PronType=Ind",
"Gender=Neut|Number=Sing|POS=PRON|PronType=Ind",
"Mood=Imp|POS=VERB",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=DET|Person=1|Poss=Yes|PronType=Prs",
"Definite=Ind|Number=Sing|POS=AUX|Tense=Past|VerbForm=Part",
"POS=X",
"Case=Nom|Gender=Com|Number=Plur|POS=PRON|Person=1|PronType=Prs",
"Case=Gen|Definite=Def|Gender=Com|Number=Plur|POS=NOUN",
"POS=VERB|Tense=Pres|VerbForm=Part",
"Number=Plur|POS=PRON|PronType=Int,Rel",
"POS=VERB|VerbForm=Inf|Voice=Pass",
"Case=Gen|Definite=Ind|Gender=Com|Number=Sing|POS=NOUN",
"Degree=Cmp|POS=ADV",
"POS=ADV|PartType=Inf",
"Degree=Sup|POS=ADV",
"Number=Plur|POS=PRON|PronType=Dem",
"Number=Plur|POS=PRON|PronType=Ind",
"Definite=Def|Gender=Neut|Number=Plur|POS=NOUN",
"Case=Acc|Gender=Com|Number=Sing|POS=PRON|Person=1|PronType=Prs",
"Case=Gen|POS=PROPN",
"POS=ADP",
"Degree=Cmp|Number=Plur|POS=ADJ",
"Definite=Def|Degree=Sup|POS=ADJ",
"Gender=Neut|Number=Sing|Number[psor]=Sing|POS=DET|Person=1|Poss=Yes|PronType=Prs",
"Degree=Pos|Number=Sing|POS=ADJ",
"Number=Plur|Number[psor]=Sing|POS=DET|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Gender=Com|Number=Sing|Number[psor]=Plur|POS=DET|Person=1|Poss=Yes|PronType=Prs|Style=Form",
"Number=Plur|POS=PRON|PronType=Rcp",
"Case=Gen|Degree=Cmp|POS=ADJ",
"POS=SPACE",
"Case=Gen|Definite=Def|Gender=Neut|Number=Plur|POS=NOUN",
"Number[psor]=Plur|POS=DET|Person=3|Poss=Yes|PronType=Prs",
"POS=INTJ",
"Number=Plur|Number[psor]=Sing|POS=DET|Person=1|Poss=Yes|PronType=Prs",
"Degree=Pos|Gender=Neut|Number=Sing|POS=ADJ",
"Gender=Neut|Number=Sing|Number[psor]=Plur|POS=DET|Person=1|Poss=Yes|PronType=Prs|Style=Form",
"Case=Acc|Gender=Com|Number=Sing|POS=PRON|Person=2|PronType=Prs",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=DET|Person=2|Poss=Yes|PronType=Prs",
"Case=Gen|Definite=Ind|Gender=Neut|Number=Plur|POS=NOUN",
"Number=Sing|POS=PRON|PronType=Int,Rel",
"Number=Plur|Number[psor]=Plur|POS=DET|Person=1|Poss=Yes|PronType=Prs|Style=Form",
"Gender=Neut|Number=Sing|POS=PRON|PronType=Int,Rel",
"Definite=Def|Degree=Sup|Number=Plur|POS=ADJ",
"Case=Nom|Gender=Com|Number=Sing|POS=PRON|Person=2|PronType=Prs",
"Gender=Neut|Number=Sing|Number[psor]=Sing|POS=DET|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Definite=Ind|Number=Sing|POS=NOUN",
"Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Plur|Number[psor]=Sing|POS=PRON|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"POS=SYM",
"Case=Nom|Gender=Com|POS=PRON|Person=2|Polite=Form|PronType=Prs",
"Degree=Sup|POS=ADJ",
"Number=Plur|POS=DET|PronType=Ind|Style=Arch",
"Case=Gen|Gender=Com|Number=Sing|POS=DET|PronType=Dem",
"Foreign=Yes|POS=X",
"POS=DET|Person=2|Polite=Form|Poss=Yes|PronType=Prs",
"Gender=Neut|Number=Sing|POS=PRON|PronType=Dem",
"Case=Acc|Gender=Com|Number=Plur|POS=PRON|Person=1|PronType=Prs",
"Case=Gen|Definite=Ind|Gender=Neut|Number=Sing|POS=NOUN",
"Case=Gen|POS=PRON|PronType=Int,Rel",
"Gender=Com|Number=Sing|POS=PRON|PronType=Dem",
"Abbr=Yes|POS=X",
"Case=Gen|Definite=Ind|Gender=Com|Number=Plur|POS=NOUN",
"Definite=Def|Degree=Abs|POS=ADJ",
"Definite=Ind|Degree=Sup|Number=Sing|POS=ADJ",
"Definite=Ind|POS=NOUN",
"Gender=Com|Number=Plur|POS=NOUN",
"Number[psor]=Plur|POS=DET|Person=1|Poss=Yes|PronType=Prs",
"Gender=Com|POS=PRON|PronType=Int,Rel",
"Case=Nom|Gender=Com|Number=Plur|POS=PRON|Person=2|PronType=Prs",
"Degree=Abs|POS=ADV",
"POS=VERB|VerbForm=Ger",
"POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Def|Degree=Sup|Number=Sing|POS=ADJ",
"Number=Plur|Number[psor]=Plur|POS=PRON|Person=1|Poss=Yes|PronType=Prs|Style=Form",
"Case=Gen|Definite=Def|Degree=Pos|Number=Sing|POS=ADJ",
"Case=Gen|Degree=Pos|Number=Plur|POS=ADJ",
"Case=Acc|Gender=Com|POS=PRON|Person=2|Polite=Form|PronType=Prs",
"Gender=Com|Number=Sing|POS=PRON|PronType=Int,Rel",
"POS=VERB|Tense=Pres",
"Case=Gen|Number=Plur|POS=DET|PronType=Ind",
"Number[psor]=Plur|POS=DET|Person=2|Poss=Yes|PronType=Prs",
"POS=PRON|Person=2|Polite=Form|Poss=Yes|PronType=Prs",
"Gender=Neut|Number=Sing|Number[psor]=Sing|POS=DET|Person=2|Poss=Yes|PronType=Prs",
"POS=AUX|Tense=Pres|VerbForm=Part",
"Mood=Ind|POS=VERB|Tense=Past|VerbForm=Fin|Voice=Pass",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=PRON|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Degree=Sup|Number=Plur|POS=ADJ",
"Case=Acc|Gender=Com|Number=Plur|POS=PRON|Person=2|PronType=Prs",
"Gender=Neut|Number=Sing|Number[psor]=Sing|POS=PRON|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Definite=Ind|Number=Plur|POS=NOUN",
"Case=Gen|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Mood=Imp|POS=AUX",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=PRON|Person=1|Poss=Yes|PronType=Prs",
"Number[psor]=Sing|POS=PRON|Person=3|Poss=Yes|PronType=Prs",
"Definite=Def|Gender=Com|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Plur|Number[psor]=Sing|POS=DET|Person=2|Poss=Yes|PronType=Prs",
"Case=Gen|Gender=Com|Number=Sing|POS=DET|PronType=Ind",
"Case=Gen|POS=NOUN",
"Number[psor]=Plur|POS=PRON|Person=3|Poss=Yes|PronType=Prs",
"POS=DET|PronType=Dem",
"Definite=Def|Number=Plur|POS=NOUN"
],
"parser":[
"ROOT",
"acl:relcl",
"advcl",
"advmod",
"advmod:lmod",
"amod",
"appos",
"aux",
"case",
"cc",
"ccomp",
"compound:prt",
"conj",
"cop",
"dep",
"det",
"expl",
"fixed",
"flat",
"iobj",
"list",
"mark",
"nmod",
"nmod:poss",
"nsubj",
"nummod",
"obj",
"obl",
"obl:lmod",
"obl:tmod",
"punct",
"xcomp"
],
"attribute_ruler":[
],
"ner":[
"LOC",
"MISC",
"ORG",
"PER"
]
},
"pipeline":[
"tok2vec",
"morphologizer",
"parser",
"lemmatizer",
"attribute_ruler",
"ner"
],
"components":[
"tok2vec",
"morphologizer",
"parser",
"lemmatizer",
"senter",
"attribute_ruler",
"ner"
],
"disabled":[
"senter"
],
"performance":{
"token_acc":0.9989350373,
"token_p":0.9977732598,
"token_r":0.9974835463,
"token_f":0.997628382,
"pos_acc":0.9665859564,
"morph_acc":0.9573849879,
"morph_micro_p":0.9742794693,
"morph_micro_r":0.967492807,
"morph_micro_f":0.9708742782,
"morph_per_feat":{
"Mood":{
"p":0.982791587,
"r":0.9799809342,
"f":0.9813842482
},
"Tense":{
"p":0.9796072508,
"r":0.9766566265,
"f":0.9781297134
},
"VerbForm":{
"p":0.9710412816,
"r":0.964504284,
"f":0.9677617439
},
"Voice":{
"p":0.983495874,
"r":0.9798206278,
"f":0.9816548109
},
"Definite":{
"p":0.9669585987,
"r":0.9596997234,
"f":0.9633154868
},
"Gender":{
"p":0.9589315526,
"r":0.9544699236,
"f":0.9566955363
},
"Number":{
"p":0.967648606,
"r":0.9595722483,
"f":0.9635935045
},
"AdpType":{
"p":1.0,
"r":0.9893899204,
"f":0.9946666667
},
"PartType":{
"p":1.0,
"r":1.0,
"f":1.0
},
"Case":{
"p":0.9856,
"r":0.9731437599,
"f":0.9793322734
},
"Person":{
"p":0.9858156028,
"r":0.9875666075,
"f":0.9866903283
},
"PronType":{
"p":0.9876441516,
"r":0.9860197368,
"f":0.9868312757
},
"NumType":{
"p":0.9863945578,
"r":0.9602649007,
"f":0.9731543624
},
"Degree":{
"p":0.9657701711,
"r":0.9518072289,
"f":0.9587378641
},
"Reflex":{
"p":1.0,
"r":1.0,
"f":1.0
},
"Number[psor]":{
"p":0.988372093,
"r":0.988372093,
"f":0.988372093
},
"Poss":{
"p":1.0,
"r":0.9886363636,
"f":0.9942857143
},
"Foreign":{
"p":1.0,
"r":0.5,
"f":0.6666666667
},
"Abbr":{
"p":0.0,
"r":0.0,
"f":0.0
},
"Style":{
"p":1.0,
"r":1.0,
"f":1.0
},
"Polite":{
"p":1.0,
"r":0.5,
"f":0.6666666667
}
},
"sents_p":0.8908765653,
"sents_r":0.8829787234,
"sents_f":0.8869100623,
"dep_uas":0.8225238813,
"dep_las":0.7828612927,
"dep_las_per_type":{
"advmod":{
"p":0.6876675603,
"r":0.7245762712,
"f":0.7056396149
},
"root":{
"p":0.8240574506,
"r":0.8138297872,
"f":0.818911686
},
"nsubj":{
"p":0.8513800425,
"r":0.8459915612,
"f":0.8486772487
},
"case":{
"p":0.8941641939,
"r":0.8915187377,
"f":0.8928395062
},
"obl":{
"p":0.7017828201,
"r":0.6723602484,
"f":0.6867565424
},
"cc":{
"p":0.795389049,
"r":0.8023255814,
"f":0.7988422576
},
"conj":{
"p":0.5918918919,
"r":0.584,
"f":0.5879194631
},
"obj":{
"p":0.7781690141,
"r":0.8582524272,
"f":0.8162511542
},
"aux":{
"p":0.8922155689,
"r":0.8688046647,
"f":0.8803545052
},
"acl:relcl":{
"p":0.606741573,
"r":0.5837837838,
"f":0.5950413223
},
"advmod:lmod":{
"p":0.7627118644,
"r":0.671641791,
"f":0.7142857143
},
"det":{
"p":0.9247135843,
"r":0.9308072488,
"f":0.9277504105
},
"amod":{
"p":0.8291032149,
"r":0.8361774744,
"f":0.8326253186
},
"nmod:poss":{
"p":0.7052631579,
"r":0.6633663366,
"f":0.6836734694
},
"ccomp":{
"p":0.5555555556,
"r":0.6451612903,
"f":0.5970149254
},
"nummod":{
"p":0.811023622,
"r":0.8583333333,
"f":0.8340080972
},
"flat":{
"p":0.7743902439,
"r":0.8410596026,
"f":0.8063492063
},
"compound:prt":{
"p":0.3888888889,
"r":0.3414634146,
"f":0.3636363636
},
"advcl":{
"p":0.6635514019,
"r":0.6120689655,
"f":0.6367713004
},
"mark":{
"p":0.8902953586,
"r":0.8665297741,
"f":0.878251821
},
"cop":{
"p":0.8222222222,
"r":0.8457142857,
"f":0.8338028169
},
"dep":{
"p":0.1111111111,
"r":0.1509433962,
"f":0.128
},
"nmod":{
"p":0.6686626747,
"r":0.654296875,
"f":0.6614017769
},
"iobj":{
"p":0.9,
"r":0.4090909091,
"f":0.5625
},
"xcomp":{
"p":0.4468085106,
"r":0.3559322034,
"f":0.3962264151
},
"list":{
"p":0.5,
"r":0.2222222222,
"f":0.3076923077
},
"vocative":{
"p":0.0,
"r":0.0,
"f":0.0
},
"fixed":{
"p":0.8888888889,
"r":0.7804878049,
"f":0.8311688312
},
"expl":{
"p":0.8529411765,
"r":0.8529411765,
"f":0.8529411765
},
"appos":{
"p":0.5862068966,
"r":0.5151515152,
"f":0.5483870968
},
"obl:tmod":{
"p":0.75,
"r":0.3333333333,
"f":0.4615384615
},
"discourse":{
"p":0.0,
"r":0.0,
"f":0.0
},
"obl:lmod":{
"p":0.0,
"r":0.0,
"f":0.0
}
},
"lemma_acc":0.948377724,
"tag_acc":0.9665859564,
"ents_p":0.800407332,
"ents_r":0.81875,
"ents_f":0.8094747683,
"ents_per_type":{
"PER":{
"p":0.893081761,
"r":0.8554216867,
"f":0.8738461538
},
"ORG":{
"p":0.7222222222,
"r":0.7222222222,
"f":0.7222222222
},
"MISC":{
"p":0.6771653543,
"r":0.7610619469,
"f":0.7166666667
},
"LOC":{
"p":0.8695652174,
"r":0.9009009009,
"f":0.8849557522
}
},
"speed":8899.0748579849
},
"sources":[
{
"name":"UD Danish DDT v2.8",
"url":"https://github.com/UniversalDependencies/UD_Danish-DDT",
"license":"CC BY-SA 4.0",
"author":"Johannsen, Anders; Mart\u00ednez Alonso, H\u00e9ctor; Plank, Barbara"
},
{
"name":"DaNE",
"url":"https://github.com/alexandrainst/danlp/blob/master/docs/datasets.md#danish-dependency-treebank-dane",
"license":"CC BY-SA 4.0",
"author":"Rasmus Hvingelby, Amalie B. Pauli, Maria Barrett, Christina Rosted, Lasse M. Lidegaard, Anders S\u00f8gaard"
},
{
"name":"Explosion fastText Vectors (cbow, OSCAR Common Crawl + Wikipedia)",
"url":"https://spacy.io",
"license":"CC0",
"author":"Explosion"
}
],
"requirements":[
]
}