da_core_news_trf / meta.json
osanseviero's picture
Update spaCy pipeline
c79eab0
raw
history blame
17.3 kB
{
"lang":"da",
"name":"core_news_trf",
"version":"3.2.0",
"description":"Danish transformer pipeline (Maltehb/danish-bert-botxo). Components: transformer, morphologizer, parser, ner, attribute_ruler, lemmatizer.",
"author":"Explosion",
"email":"[email protected]",
"url":"https://explosion.ai",
"license":"CC BY-SA 4.0",
"spacy_version":">=3.2.0,<3.3.0",
"spacy_git_version":"bb26550e2",
"vectors":{
"width":0,
"vectors":0,
"keys":0,
"name":null
},
"labels":{
"transformer":[
],
"morphologizer":[
"AdpType=Prep|POS=ADP",
"Definite=Ind|Gender=Com|Number=Sing|POS=NOUN",
"Mood=Ind|POS=AUX|Tense=Pres|VerbForm=Fin|Voice=Act",
"POS=PROPN",
"Definite=Ind|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Def|Gender=Neut|Number=Sing|POS=NOUN",
"POS=SCONJ",
"Definite=Def|Gender=Com|Number=Sing|POS=NOUN",
"Mood=Ind|POS=VERB|Tense=Pres|VerbForm=Fin|Voice=Act",
"POS=ADV",
"Number=Plur|POS=DET|PronType=Dem",
"Degree=Pos|Number=Plur|POS=ADJ",
"Definite=Ind|Gender=Com|Number=Plur|POS=NOUN",
"POS=PUNCT",
"POS=CCONJ",
"Definite=Ind|Degree=Cmp|Number=Sing|POS=ADJ",
"Degree=Cmp|POS=ADJ",
"POS=PRON|PartType=Inf",
"Gender=Com|Number=Sing|POS=DET|PronType=Ind",
"Definite=Ind|Degree=Pos|Number=Sing|POS=ADJ",
"Case=Acc|Gender=Neut|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Definite=Ind|Gender=Neut|Number=Plur|POS=NOUN",
"Definite=Def|Degree=Pos|Number=Sing|POS=ADJ",
"Gender=Neut|Number=Sing|POS=DET|PronType=Dem",
"Degree=Pos|POS=ADV",
"Definite=Def|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Ind|Gender=Neut|Number=Sing|POS=NOUN",
"POS=PRON|PronType=Dem",
"NumType=Card|POS=NUM",
"Definite=Ind|Degree=Pos|Gender=Neut|Number=Sing|POS=ADJ",
"Case=Acc|Gender=Com|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"Degree=Pos|Gender=Com|Number=Sing|POS=ADJ",
"Case=Nom|Gender=Com|Number=Sing|POS=PRON|Person=3|PronType=Prs",
"NumType=Ord|POS=ADJ",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=DET|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Mood=Ind|POS=AUX|Tense=Past|VerbForm=Fin|Voice=Act",
"POS=VERB|VerbForm=Inf|Voice=Act",
"Mood=Ind|POS=VERB|Tense=Past|VerbForm=Fin|Voice=Act",
"POS=NOUN",
"Mood=Ind|POS=VERB|Tense=Pres|VerbForm=Fin|Voice=Pass",
"POS=ADP|PartType=Inf",
"Degree=Pos|POS=ADJ",
"Definite=Def|Gender=Com|Number=Plur|POS=NOUN",
"Number[psor]=Sing|POS=DET|Person=3|Poss=Yes|PronType=Prs",
"Case=Gen|Definite=Def|Gender=Com|Number=Sing|POS=NOUN",
"POS=AUX|VerbForm=Inf|Voice=Act",
"Definite=Ind|Degree=Pos|Gender=Com|Number=Sing|POS=ADJ",
"Gender=Com|Number=Sing|POS=DET|PronType=Dem",
"Number=Plur|POS=DET|PronType=Ind",
"Gender=Com|Number=Sing|POS=PRON|PronType=Ind",
"Case=Acc|POS=PRON|Person=3|PronType=Prs|Reflex=Yes",
"POS=PART|PartType=Inf",
"Gender=Neut|Number=Sing|POS=DET|PronType=Ind",
"Case=Acc|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Case=Gen|Definite=Def|Gender=Neut|Number=Sing|POS=NOUN",
"Case=Nom|Number=Plur|POS=PRON|Person=3|PronType=Prs",
"Case=Nom|Gender=Com|Number=Sing|POS=PRON|Person=1|PronType=Prs",
"Case=Nom|Gender=Com|POS=PRON|PronType=Ind",
"Gender=Neut|Number=Sing|POS=PRON|PronType=Ind",
"Mood=Imp|POS=VERB",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=DET|Person=1|Poss=Yes|PronType=Prs",
"Definite=Ind|Number=Sing|POS=AUX|Tense=Past|VerbForm=Part",
"POS=X",
"Case=Nom|Gender=Com|Number=Plur|POS=PRON|Person=1|PronType=Prs",
"Case=Gen|Definite=Def|Gender=Com|Number=Plur|POS=NOUN",
"POS=VERB|Tense=Pres|VerbForm=Part",
"Number=Plur|POS=PRON|PronType=Int,Rel",
"POS=VERB|VerbForm=Inf|Voice=Pass",
"Case=Gen|Definite=Ind|Gender=Com|Number=Sing|POS=NOUN",
"Degree=Cmp|POS=ADV",
"POS=ADV|PartType=Inf",
"Degree=Sup|POS=ADV",
"Number=Plur|POS=PRON|PronType=Dem",
"Number=Plur|POS=PRON|PronType=Ind",
"Definite=Def|Gender=Neut|Number=Plur|POS=NOUN",
"Case=Acc|Gender=Com|Number=Sing|POS=PRON|Person=1|PronType=Prs",
"Case=Gen|POS=PROPN",
"POS=ADP",
"Degree=Cmp|Number=Plur|POS=ADJ",
"Definite=Def|Degree=Sup|POS=ADJ",
"Gender=Neut|Number=Sing|Number[psor]=Sing|POS=DET|Person=1|Poss=Yes|PronType=Prs",
"Degree=Pos|Number=Sing|POS=ADJ",
"Number=Plur|Number[psor]=Sing|POS=DET|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Gender=Com|Number=Sing|Number[psor]=Plur|POS=DET|Person=1|Poss=Yes|PronType=Prs|Style=Form",
"Number=Plur|POS=PRON|PronType=Rcp",
"Case=Gen|Degree=Cmp|POS=ADJ",
"Case=Gen|Definite=Def|Gender=Neut|Number=Plur|POS=NOUN",
"Number[psor]=Plur|POS=DET|Person=3|Poss=Yes|PronType=Prs",
"POS=INTJ",
"Number=Plur|Number[psor]=Sing|POS=DET|Person=1|Poss=Yes|PronType=Prs",
"Degree=Pos|Gender=Neut|Number=Sing|POS=ADJ",
"Gender=Neut|Number=Sing|Number[psor]=Plur|POS=DET|Person=1|Poss=Yes|PronType=Prs|Style=Form",
"Case=Acc|Gender=Com|Number=Sing|POS=PRON|Person=2|PronType=Prs",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=DET|Person=2|Poss=Yes|PronType=Prs",
"Case=Gen|Definite=Ind|Gender=Neut|Number=Plur|POS=NOUN",
"Number=Sing|POS=PRON|PronType=Int,Rel",
"Number=Plur|Number[psor]=Plur|POS=DET|Person=1|Poss=Yes|PronType=Prs|Style=Form",
"Gender=Neut|Number=Sing|POS=PRON|PronType=Int,Rel",
"Definite=Def|Degree=Sup|Number=Plur|POS=ADJ",
"Case=Nom|Gender=Com|Number=Sing|POS=PRON|Person=2|PronType=Prs",
"Gender=Neut|Number=Sing|Number[psor]=Sing|POS=DET|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Definite=Ind|Number=Sing|POS=NOUN",
"Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Plur|Number[psor]=Sing|POS=PRON|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"POS=SYM",
"Case=Nom|Gender=Com|POS=PRON|Person=2|Polite=Form|PronType=Prs",
"Degree=Sup|POS=ADJ",
"Number=Plur|POS=DET|PronType=Ind|Style=Arch",
"Case=Gen|Gender=Com|Number=Sing|POS=DET|PronType=Dem",
"Foreign=Yes|POS=X",
"POS=DET|Person=2|Polite=Form|Poss=Yes|PronType=Prs",
"Gender=Neut|Number=Sing|POS=PRON|PronType=Dem",
"Case=Acc|Gender=Com|Number=Plur|POS=PRON|Person=1|PronType=Prs",
"Case=Gen|Definite=Ind|Gender=Neut|Number=Sing|POS=NOUN",
"Case=Gen|POS=PRON|PronType=Int,Rel",
"Gender=Com|Number=Sing|POS=PRON|PronType=Dem",
"Abbr=Yes|POS=X",
"Case=Gen|Definite=Ind|Gender=Com|Number=Plur|POS=NOUN",
"Definite=Def|Degree=Abs|POS=ADJ",
"Definite=Ind|Degree=Sup|Number=Sing|POS=ADJ",
"Definite=Ind|POS=NOUN",
"Gender=Com|Number=Plur|POS=NOUN",
"Number[psor]=Plur|POS=DET|Person=1|Poss=Yes|PronType=Prs",
"Gender=Com|POS=PRON|PronType=Int,Rel",
"Case=Nom|Gender=Com|Number=Plur|POS=PRON|Person=2|PronType=Prs",
"Degree=Abs|POS=ADV",
"POS=VERB|VerbForm=Ger",
"POS=VERB|Tense=Past|VerbForm=Part",
"Definite=Def|Degree=Sup|Number=Sing|POS=ADJ",
"Number=Plur|Number[psor]=Plur|POS=PRON|Person=1|Poss=Yes|PronType=Prs|Style=Form",
"Case=Gen|Definite=Def|Degree=Pos|Number=Sing|POS=ADJ",
"Case=Gen|Degree=Pos|Number=Plur|POS=ADJ",
"Case=Acc|Gender=Com|POS=PRON|Person=2|Polite=Form|PronType=Prs",
"Gender=Com|Number=Sing|POS=PRON|PronType=Int,Rel",
"POS=VERB|Tense=Pres",
"Case=Gen|Number=Plur|POS=DET|PronType=Ind",
"Number[psor]=Plur|POS=DET|Person=2|Poss=Yes|PronType=Prs",
"POS=PRON|Person=2|Polite=Form|Poss=Yes|PronType=Prs",
"Gender=Neut|Number=Sing|Number[psor]=Sing|POS=DET|Person=2|Poss=Yes|PronType=Prs",
"POS=AUX|Tense=Pres|VerbForm=Part",
"Mood=Ind|POS=VERB|Tense=Past|VerbForm=Fin|Voice=Pass",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=PRON|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Degree=Sup|Number=Plur|POS=ADJ",
"Case=Acc|Gender=Com|Number=Plur|POS=PRON|Person=2|PronType=Prs",
"Gender=Neut|Number=Sing|Number[psor]=Sing|POS=PRON|Person=3|Poss=Yes|PronType=Prs|Reflex=Yes",
"Definite=Ind|Number=Plur|POS=NOUN",
"Case=Gen|Number=Plur|POS=VERB|Tense=Past|VerbForm=Part",
"Mood=Imp|POS=AUX",
"Gender=Com|Number=Sing|Number[psor]=Sing|POS=PRON|Person=1|Poss=Yes|PronType=Prs",
"Number[psor]=Sing|POS=PRON|Person=3|Poss=Yes|PronType=Prs",
"Definite=Def|Gender=Com|Number=Sing|POS=VERB|Tense=Past|VerbForm=Part",
"Number=Plur|Number[psor]=Sing|POS=DET|Person=2|Poss=Yes|PronType=Prs",
"Case=Gen|Gender=Com|Number=Sing|POS=DET|PronType=Ind",
"Case=Gen|POS=NOUN",
"Number[psor]=Plur|POS=PRON|Person=3|Poss=Yes|PronType=Prs",
"POS=DET|PronType=Dem",
"Definite=Def|Number=Plur|POS=NOUN"
],
"parser":[
"ROOT",
"acl:relcl",
"advcl",
"advmod",
"advmod:lmod",
"amod",
"appos",
"aux",
"case",
"cc",
"ccomp",
"compound:prt",
"conj",
"cop",
"dep",
"det",
"expl",
"fixed",
"flat",
"iobj",
"list",
"mark",
"nmod",
"nmod:poss",
"nsubj",
"nummod",
"obj",
"obl",
"obl:lmod",
"obl:tmod",
"punct",
"xcomp"
],
"attribute_ruler":[
],
"lemmatizer":[
],
"ner":[
"LOC",
"MISC",
"ORG",
"PER"
]
},
"pipeline":[
"transformer",
"morphologizer",
"parser",
"attribute_ruler",
"lemmatizer",
"ner"
],
"components":[
"transformer",
"morphologizer",
"parser",
"attribute_ruler",
"lemmatizer",
"ner"
],
"disabled":[
],
"performance":{
"token_acc":0.9994672349,
"token_p":0.9977732598,
"token_r":0.9974835463,
"token_f":0.997628382,
"pos_acc":0.9734127561,
"morph_acc":0.9700227614,
"morph_micro_p":0.9874419317,
"morph_micro_r":0.9763767604,
"morph_micro_f":0.9818781726,
"morph_per_feat":{
"Mood":{
"p":0.9923298178,
"r":0.9866539561,
"f":0.9894837476
},
"Tense":{
"p":0.9832953683,
"r":0.9751506024,
"f":0.9792060491
},
"VerbForm":{
"p":0.9833127318,
"r":0.9736842105,
"f":0.9784747847
},
"Voice":{
"p":0.9932381668,
"r":0.9880418535,
"f":0.990633196
},
"Definite":{
"p":0.9899678973,
"r":0.974713552,
"f":0.9822815051
},
"Gender":{
"p":0.9814690027,
"r":0.9680957129,
"f":0.9747364899
},
"Number":{
"p":0.9889006342,
"r":0.9760041732,
"f":0.9824100814
},
"AdpType":{
"p":0.9946428571,
"r":0.9849690539,
"f":0.989782319
},
"PartType":{
"p":1.0,
"r":0.9967532468,
"f":0.9983739837
},
"Case":{
"p":0.9935794543,
"r":0.9778830964,
"f":0.9856687898
},
"Person":{
"p":0.9857397504,
"r":0.9822380107,
"f":0.9839857651
},
"PronType":{
"p":0.9950166113,
"r":0.9851973684,
"f":0.9900826446
},
"NumType":{
"p":0.9731543624,
"r":0.9602649007,
"f":0.9666666667
},
"Degree":{
"p":0.9696233293,
"r":0.9614457831,
"f":0.9655172414
},
"Reflex":{
"p":1.0,
"r":1.0,
"f":1.0
},
"Polite":{
"p":0.0,
"r":0.0,
"f":0.0
},
"Number[psor]":{
"p":0.9770114943,
"r":0.988372093,
"f":0.9826589595
},
"Poss":{
"p":1.0,
"r":0.9886363636,
"f":0.9942857143
},
"Foreign":{
"p":0.875,
"r":0.7,
"f":0.7777777778
},
"Abbr":{
"p":1.0,
"r":0.4,
"f":0.5714285714
},
"Style":{
"p":1.0,
"r":1.0,
"f":1.0
}
},
"sents_p":0.8057921635,
"sents_r":0.8386524823,
"sents_f":0.8218940052,
"dep_uas":0.8639535533,
"dep_las":0.8326913415,
"dep_las_per_type":{
"advmod":{
"p":0.7794729542,
"r":0.7937853107,
"f":0.7865640308
},
"root":{
"p":0.8398637138,
"r":0.8741134752,
"f":0.8566463944
},
"nsubj":{
"p":0.9073482428,
"r":0.8987341772,
"f":0.9030206677
},
"case":{
"p":0.9097291876,
"r":0.8944773176,
"f":0.9020387867
},
"obl":{
"p":0.7601809955,
"r":0.7826086957,
"f":0.7712318286
},
"cc":{
"p":0.8397626113,
"r":0.8226744186,
"f":0.8311306902
},
"conj":{
"p":0.7506849315,
"r":0.7306666667,
"f":0.7405405405
},
"obj":{
"p":0.868852459,
"r":0.9262135922,
"f":0.8966165414
},
"aux":{
"p":0.8941176471,
"r":0.8862973761,
"f":0.8901903367
},
"acl:relcl":{
"p":0.7719298246,
"r":0.7135135135,
"f":0.7415730337
},
"advmod:lmod":{
"p":0.7846153846,
"r":0.7611940299,
"f":0.7727272727
},
"det":{
"p":0.9048387097,
"r":0.9242174629,
"f":0.9144254279
},
"amod":{
"p":0.8640275387,
"r":0.8566552901,
"f":0.8603256213
},
"nmod:poss":{
"p":0.7087378641,
"r":0.7227722772,
"f":0.7156862745
},
"ccomp":{
"p":0.7419354839,
"r":0.7419354839,
"f":0.7419354839
},
"nummod":{
"p":0.8196721311,
"r":0.8333333333,
"f":0.826446281
},
"flat":{
"p":0.8291139241,
"r":0.8675496689,
"f":0.8478964401
},
"compound:prt":{
"p":0.575,
"r":0.5609756098,
"f":0.5679012346
},
"advcl":{
"p":0.664,
"r":0.7155172414,
"f":0.6887966805
},
"mark":{
"p":0.9201680672,
"r":0.8993839836,
"f":0.9096573209
},
"cop":{
"p":0.8833333333,
"r":0.9085714286,
"f":0.8957746479
},
"dep":{
"p":0.1844660194,
"r":0.358490566,
"f":0.2435897436
},
"nmod":{
"p":0.7558386412,
"r":0.6953125,
"f":0.7243133266
},
"iobj":{
"p":0.875,
"r":0.6363636364,
"f":0.7368421053
},
"obl:lmod":{
"p":0.0,
"r":0.0,
"f":0.0
},
"xcomp":{
"p":0.5714285714,
"r":0.3389830508,
"f":0.4255319149
},
"list":{
"p":0.25,
"r":0.1111111111,
"f":0.1538461538
},
"vocative":{
"p":0.0,
"r":0.0,
"f":0.0
},
"fixed":{
"p":0.9428571429,
"r":0.8048780488,
"f":0.8684210526
},
"expl":{
"p":0.96875,
"r":0.9117647059,
"f":0.9393939394
},
"appos":{
"p":0.6551724138,
"r":0.5757575758,
"f":0.6129032258
},
"obl:tmod":{
"p":0.8181818182,
"r":0.5,
"f":0.6206896552
},
"discourse":{
"p":0.0,
"r":0.0,
"f":0.0
}
},
"tag_acc":0.9734127561,
"lemma_acc":0.8491041162,
"ents_p":0.7784313725,
"ents_r":0.8270833333,
"ents_f":0.802020202,
"ents_per_type":{
"PER":{
"p":0.8982035928,
"r":0.9036144578,
"f":0.9009009009
},
"ORG":{
"p":0.7058823529,
"r":0.6666666667,
"f":0.6857142857
},
"MISC":{
"p":0.6769230769,
"r":0.7787610619,
"f":0.7242798354
},
"LOC":{
"p":0.7734375,
"r":0.8918918919,
"f":0.8284518828
}
},
"speed":5370.7462201589
},
"sources":[
{
"name":"UD Danish DDT v2.8",
"url":"https://github.com/UniversalDependencies/UD_Danish-DDT",
"license":"CC BY-SA 4.0",
"author":"Johannsen, Anders; Mart\u00ednez Alonso, H\u00e9ctor; Plank, Barbara"
},
{
"name":"DaNE",
"url":"https://github.com/alexandrainst/danlp/blob/master/docs/datasets.md#danish-dependency-treebank-dane",
"license":"CC BY-SA 4.0",
"author":"Rasmus Hvingelby, Amalie B. Pauli, Maria Barrett, Christina Rosted, Lasse M. Lidegaard, Anders S\u00f8gaard"
},
{
"name":"Sprogteknologisk orddatabase over det danske sprog",
"url":"https://cst.ku.dk/sto_ordbase/",
"license":"CC BY-SA 4.0",
"author":"Center for Language Technology, University of Copenhagen"
},
{
"name":"Maltehb/danish-bert-botxo",
"author":"BotXO.ai",
"url":"https://huggingface.co./Maltehb/danish-bert-botxo",
"license":"CC BY 4.0"
}
],
"requirements":[
"spacy-transformers>=1.1.2,<1.2.0"
]
}