Add evaluation results on the autoevaluate--xsum-sample config and test split of autoevaluate/xsum-sample

#7
by autoevaluator HF staff - opened
Files changed (1) hide show
  1. README.md +42 -3
README.md CHANGED
@@ -12,16 +12,55 @@ model-index:
12
  - name: summarization
13
  results:
14
  - task:
15
- name: Sequence-to-sequence Language Modeling
16
  type: text2text-generation
 
17
  dataset:
18
  name: xsum
19
  type: xsum
20
  args: default
21
  metrics:
22
- - name: Rouge1
23
- type: rouge
24
  value: 23.9405
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
25
  ---
26
 
27
  <!-- This model card has been generated automatically according to the information the Trainer had access to. You
 
12
  - name: summarization
13
  results:
14
  - task:
 
15
  type: text2text-generation
16
+ name: Sequence-to-sequence Language Modeling
17
  dataset:
18
  name: xsum
19
  type: xsum
20
  args: default
21
  metrics:
22
+ - type: rouge
 
23
  value: 23.9405
24
+ name: Rouge1
25
+ - task:
26
+ type: summarization
27
+ name: Summarization
28
+ dataset:
29
+ name: autoevaluate/xsum-sample
30
+ type: autoevaluate/xsum-sample
31
+ config: autoevaluate--xsum-sample
32
+ split: test
33
+ metrics:
34
+ - type: rouge
35
+ value: 18.3598
36
+ name: ROUGE-1
37
+ verified: true
38
+ verifyToken: eyJhbGciOiJFZERTQSIsInR5cCI6IkpXVCJ9.eyJoYXNoIjoiMGJiZThmYzcwMDU4MjI4MjZlNTBjNDQ5MjYxZDNiNzU3Y2Y5OWMxZmZjODAwYTI0YTdkZTZmZTVjMmI3MGY0MSIsInZlcnNpb24iOjF9.KscbYOZebwlZfpNk-X0c54yQc7T4sXDa-ICk3WMLBsDFYAT4RGeSOa7YZTbBKaQ9ebMgl9adQ0PPV4u6t3vNBQ
39
+ - type: rouge
40
+ value: 3.0796
41
+ name: ROUGE-2
42
+ verified: true
43
+ verifyToken: eyJhbGciOiJFZERTQSIsInR5cCI6IkpXVCJ9.eyJoYXNoIjoiZWY2NjViZWYxODYyYWQyYzliYjVlY2NiYzU2YmYyNzQ5MTE5NTE0MjhhMTU3NDk0YjVhYjRmOGM3YmNjOGE1NyIsInZlcnNpb24iOjF9.vq-_esb4DMnlJ_lbzSY_AClBMSBnTNuzCKeVkIEc_vxqXrKI7Dz7pYkxnzGXOznXc0gTkfGGp7kOUzYDc3cdDQ
44
+ - type: rouge
45
+ value: 14.9038
46
+ name: ROUGE-L
47
+ verified: true
48
+ verifyToken: eyJhbGciOiJFZERTQSIsInR5cCI6IkpXVCJ9.eyJoYXNoIjoiMGQ1ODA0NzA1N2QzZGQ4MDgyMWE0YzQ1ZmQ5YWFhZTEzYmZhMjUwYTYyOWU5MzdjZGUxNThkZmQ4OGY0MmVkNiIsInZlcnNpb24iOjF9.BFaFzEuNfDuwqpiFjb8HY6uRQdgzg41plZuU8eEeBjHJSF1QvNwA6oWvUCSToT-LqftjYuoy_-jgNsFd-sziCA
49
+ - type: rouge
50
+ value: 14.8069
51
+ name: ROUGE-LSUM
52
+ verified: true
53
+ verifyToken: eyJhbGciOiJFZERTQSIsInR5cCI6IkpXVCJ9.eyJoYXNoIjoiOTE0MmU4MDAzZjI4ZDFiZTEzZWQwMGMzZGZmZTM1ZmM3ZTM4ZWE1OWFmZTI3NDg3OTA0ZDRlNzU5YzQ0NWI1ZSIsInZlcnNpb24iOjF9.aMSrcZ0bh_H66JSnClCOYiozFPUSa9xxKn_4xjqRGsjNX9nv6ELVeNpPJNO4w8gbZxT8RkeZJx99_t7F_7u6BA
54
+ - type: loss
55
+ value: 3.009582281112671
56
+ name: loss
57
+ verified: true
58
+ verifyToken: eyJhbGciOiJFZERTQSIsInR5cCI6IkpXVCJ9.eyJoYXNoIjoiNjJkMmRiNmMzYTAzMWVhNTJjYmM2ZjRkZDg1M2FjYTdiNmUyY2RmYjIwYjdlODQ3OTY3YjI0ZWUwNWFjNWEyZCIsInZlcnNpb24iOjF9.3z4IZp7P5WPZ3lFyjTcHVMZy2eKhlh8sp6zZno8XstvFQqt7vcSljfx1sH_9GcC8xtNL0b83r2qZKL8Zc_8gCQ
59
+ - type: gen_len
60
+ value: 18.05
61
+ name: gen_len
62
+ verified: true
63
+ verifyToken: eyJhbGciOiJFZERTQSIsInR5cCI6IkpXVCJ9.eyJoYXNoIjoiMDgzZDBjMzZmMmQ5Zjc3NDYwYzhmNGY2ZDA1ZDlkMGI5OTU2N2RkYjUzZGNiM2YwYTU4MDhiZTQxNTkxZDIyNyIsInZlcnNpb24iOjF9.FIorCL9Gpp2MoftgKvST5bj_WTjDP7KkxclK1JOiN9dTyzQDsaG1wIoUewm4NV9BMXTDJkFORi39DypL9NRZBw
64
  ---
65
 
66
  <!-- This model card has been generated automatically according to the information the Trainer had access to. You