ja_core_news_md / meta.json
osanseviero's picture
osanseviero HF staff
Update spaCy pipeline
5d3b653
raw history blame
No virus
6.99 kB
{
"lang":"ja",
"name":"core_news_md",
"version":"3.1.0",
"description":"Japanese pipeline optimized for CPU. Components: tok2vec, parser, senter, ner, attribute_ruler.",
"author":"Explosion",
"email":"contact@explosion.ai",
"url":"https://explosion.ai",
"license":"CC BY-SA 4.0",
"spacy_version":">=3.1.0,<3.2.0",
"spacy_git_version":"caba63b74",
"vectors":{
"width":300,
"vectors":20000,
"keys":480443,
"name":"ja_vectors"
},
"labels":{
"tok2vec":[
],
"parser":[
"ROOT",
"acl",
"advcl",
"advmod",
"amod",
"aux",
"case",
"cc",
"ccomp",
"compound",
"cop",
"csubj",
"dep",
"det",
"dislocated",
"fixed",
"mark",
"nmod",
"nsubj",
"nummod",
"obj",
"obl",
"punct"
],
"senter":[
"I",
"S"
],
"attribute_ruler":[
],
"ner":[
"CARDINAL",
"DATE",
"EVENT",
"FAC",
"GPE",
"LANGUAGE",
"LAW",
"LOC",
"MONEY",
"MOVEMENT",
"NORP",
"ORDINAL",
"ORG",
"PERCENT",
"PERSON",
"PET_NAME",
"PHONE",
"PRODUCT",
"QUANTITY",
"TIME",
"TITLE_AFFIX",
"WORK_OF_ART"
]
},
"pipeline":[
"tok2vec",
"parser",
"attribute_ruler",
"ner"
],
"components":[
"tok2vec",
"parser",
"senter",
"attribute_ruler",
"ner"
],
"disabled":[
"senter"
],
"performance":{
"token_acc":0.9968965945,
"tag_acc":0.9721899386,
"pos_acc":0.9639755682,
"morph_acc":0.0,
"dep_uas":0.9192898975,
"dep_las":0.9005714286,
"ents_p":0.7527932961,
"ents_r":0.6883780332,
"ents_f":0.7191460974,
"sents_p":0.994,
"sents_r":0.9920159681,
"sents_f":0.993006993,
"speed":13117.9406369597,
"morph_per_feat":{
"Polarity":{
"p":0.0,
"r":0.0,
"f":0.0
}
},
"dep_las_per_type":{
"cc":{
"p":0.829787234,
"r":0.847826087,
"f":0.8387096774
},
"nummod":{
"p":0.9806949807,
"r":0.8728522337,
"f":0.9236363636
},
"compound":{
"p":0.9418825016,
"r":0.9255121043,
"f":0.9336255479
},
"obl":{
"p":0.7983587339,
"r":0.8185096154,
"f":0.8083086053
},
"case":{
"p":0.9838709677,
"r":0.9782359679,
"f":0.9810453762
},
"dislocated":{
"p":0.5,
"r":0.3157894737,
"f":0.3870967742
},
"nmod":{
"p":0.8684546616,
"r":0.8192771084,
"f":0.843149411
},
"nsubj":{
"p":0.8020618557,
"r":0.8104166667,
"f":0.8062176166
},
"root":{
"p":0.9634888438,
"r":0.9481037924,
"f":0.9557344064
},
"aux":{
"p":0.9575844716,
"r":0.9652173913,
"f":0.9613857813
},
"advcl":{
"p":0.6797235023,
"r":0.6876456876,
"f":0.6836616454
},
"mark":{
"p":0.9575757576,
"r":0.9330708661,
"f":0.9451645065
},
"acl":{
"p":0.8216704289,
"r":0.8017621145,
"f":0.8115942029
},
"obj":{
"p":0.9320987654,
"r":0.9207317073,
"f":0.9263803681
},
"fixed":{
"p":0.9371727749,
"r":0.9835164835,
"f":0.9597855228
},
"advmod":{
"p":0.6823529412,
"r":0.4603174603,
"f":0.5497630332
},
"amod":{
"p":0.9259259259,
"r":0.625,
"f":0.7462686567
},
"cop":{
"p":0.9653179191,
"r":0.9175824176,
"f":0.9408450704
},
"ccomp":{
"p":0.95,
"r":0.8636363636,
"f":0.9047619048
},
"csubj":{
"p":0.5294117647,
"r":0.6923076923,
"f":0.6
},
"dep":{
"p":0.0,
"r":0.0,
"f":0.0
},
"det":{
"p":0.9607843137,
"r":0.9607843137,
"f":0.9607843137
}
},
"ents_per_type":{
"DATE":{
"p":0.9622641509,
"r":0.9444444444,
"f":0.953271028
},
"PERSON":{
"p":0.7769784173,
"r":0.7769784173,
"f":0.7769784173
},
"ORG":{
"p":0.6052631579,
"r":0.5267175573,
"f":0.5632653061
},
"TITLE_AFFIX":{
"p":0.84,
"r":0.7,
"f":0.7636363636
},
"GPE":{
"p":0.6813186813,
"r":0.6595744681,
"f":0.6702702703
},
"PRODUCT":{
"p":0.4285714286,
"r":0.3658536585,
"f":0.3947368421
},
"TIME":{
"p":0.6666666667,
"r":1.0,
"f":0.8
},
"QUANTITY":{
"p":0.9242424242,
"r":0.9242424242,
"f":0.9242424242
},
"NORP":{
"p":0.7777777778,
"r":0.65625,
"f":0.7118644068
},
"ORDINAL":{
"p":0.6842105263,
"r":0.6842105263,
"f":0.6842105263
},
"WORK_OF_ART":{
"p":0.8571428571,
"r":0.7058823529,
"f":0.7741935484
},
"FAC":{
"p":0.5555555556,
"r":0.4054054054,
"f":0.46875
},
"EVENT":{
"p":0.8571428571,
"r":0.4615384615,
"f":0.6
},
"PERCENT":{
"p":1.0,
"r":0.2857142857,
"f":0.4444444444
},
"LOC":{
"p":0.5625,
"r":0.9,
"f":0.6923076923
},
"MOVEMENT":{
"p":0.0,
"r":0.0,
"f":0.0
},
"LAW":{
"p":0.0,
"r":0.0,
"f":0.0
},
"MONEY":{
"p":1.0,
"r":1.0,
"f":1.0
},
"LANGUAGE":{
"p":1.0,
"r":1.0,
"f":1.0
},
"CARDINAL":{
"p":0.0,
"r":0.0,
"f":0.0
}
}
},
"sources":[
{
"name":"UD Japanese GSD v2.6",
"url":"https://github.com/UniversalDependencies/UD_Japanese-GSD",
"license":"CC BY-SA 4.0",
"author":"Omura, Mai; Miyao, Yusuke; Kanayama, Hiroshi; Matsuda, Hiroshi; Wakasa, Aya; Yamashita, Kayo; Asahara, Masayuki; Tanaka, Takaaki; Murawaki, Yugo; Matsumoto, Yuji; Mori, Shinsuke; Uematsu, Sumire; McDonald, Ryan; Nivre, Joakim; Zeman, Daniel"
},
{
"name":"UD Japanese GSD v2.6 NER",
"url":"https://github.com/megagonlabs/UD_Japanese-GSD",
"license":"CC BY-SA 4.0",
"author":"Megagon Labs Tokyo"
},
{
"name":"chiVe: Japanese Word Embedding with Sudachi & NWJC (chive-1.1-mc90-500k)",
"url":"https://github.com/WorksApplications/chiVe",
"license":"Apache-2.0",
"author":"Works Applications"
}
],
"requirements":[
"sudachipy>=0.4.9",
"sudachidict-core>=20200330"
]
}