chen459664 commited on
Commit
514ccae
·
verified ·
1 Parent(s): 1216cb5

Add files using upload-large-folder tool

Browse files
This view is limited to 50 files because it contains too many changes.   See raw diff
Files changed (50) hide show
  1. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_fra.yaml +4 -0
  2. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_ibo.yaml +4 -0
  3. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_sna.yaml +4 -0
  4. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_yaml +34 -0
  5. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_yor.yaml +4 -0
  6. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_amh.yaml +4 -0
  7. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_fra.yaml +4 -0
  8. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_ibo.yaml +4 -0
  9. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_lin.yaml +4 -0
  10. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_lug.yaml +4 -0
  11. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_sot.yaml +4 -0
  12. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_swa.yaml +4 -0
  13. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_vai.yaml +4 -0
  14. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_xho.yaml +4 -0
  15. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_yor.yaml +4 -0
  16. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_zul.yaml +4 -0
  17. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_ewe.yaml +7 -0
  18. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_hau.yaml +7 -0
  19. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_kin.yaml +7 -0
  20. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_lug.yaml +7 -0
  21. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_orm.yaml +7 -0
  22. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_sot.yaml +7 -0
  23. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_swa.yaml +7 -0
  24. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_vai.yaml +7 -0
  25. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_wol.yaml +7 -0
  26. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_yaml +33 -0
  27. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_yor.yaml +7 -0
  28. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_amh.yaml +7 -0
  29. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_eng.yaml +7 -0
  30. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_ewe.yaml +6 -0
  31. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_fra.yaml +6 -0
  32. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_hau.yaml +6 -0
  33. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_ibo.yaml +6 -0
  34. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_kin.yaml +7 -0
  35. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_orm.yaml +6 -0
  36. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_sna.yaml +7 -0
  37. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_twi.yaml +6 -0
  38. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_wol.yaml +6 -0
  39. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_xho.yaml +7 -0
  40. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_yaml +33 -0
  41. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_zul.yaml +6 -0
  42. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/afrimgsm_cot.yaml +9 -0
  43. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_amh.yaml +4 -0
  44. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_fra.yaml +4 -0
  45. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_kin.yaml +4 -0
  46. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_lug.yaml +4 -0
  47. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_orm.yaml +4 -0
  48. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_sna.yaml +4 -0
  49. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_sot.yaml +4 -0
  50. lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_swa.yaml +4 -0
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_fra.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: fra
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_fra_prompt_2
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_ibo.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: ibo
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_ibo_prompt_2
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_sna.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: sna
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_sna_prompt_2
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_yaml ADDED
@@ -0,0 +1,34 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ tag:
2
+ - afrimgsm_tasks
3
+ - afrimgsm_tasks_prompt_1
4
+ dataset_path: masakhane/afrimgsm
5
+ output_type: generate_until
6
+ test_split: test
7
+ doc_to_target: '{% if answer is not none %}{{answer[21:]}}{% else %}{{answer_number|string}}{% endif %}'
8
+ doc_to_text: "Give direct numerical answers for the question provided. \n\nQuestion: {{question}} \nAnswer: "
9
+ target_delimiter: ""
10
+ generation_kwargs:
11
+ do_sample: false
12
+ until:
13
+ - 'Question:'
14
+ - </s>
15
+ - <|im_end|>
16
+ filter_list:
17
+ - name: remove_whitespace
18
+ filter:
19
+ - function: remove_whitespace
20
+ - function: take_first
21
+ - filter:
22
+ - function: regex
23
+ group_select: -1
24
+ regex_pattern: (-?[$0-9.,]{2,})|(-?[0-9]+)
25
+ - function: take_first
26
+ name: flexible-extract
27
+ metric_list:
28
+ - metric: exact_match
29
+ aggregation: mean
30
+ higher_is_better: true
31
+ ignore_case: true
32
+ ignore_punctuation: true
33
+ metadata:
34
+ version: 2.0
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_2/afrimgsm_yor.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: yor
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_yor_prompt_2
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_amh.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: amh
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_amh_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_fra.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: fra
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_fra_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_ibo.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: ibo
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_ibo_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_lin.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: lin
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_lin_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_lug.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: lug
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_lug_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_sot.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: sot
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_sot_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_swa.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: swa
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_swa_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_vai.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: vai
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_vai_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_xho.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: xho
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_xho_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_yor.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: yor
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_yor_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_3/afrimgsm_zul.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: zul
3
+ include: afrimgsm_yaml
4
+ task: afrimgsm_zul_prompt_3
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_ewe.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: ewe
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_ewe_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_hau.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: hau
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_hau_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_kin.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: kin
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_kin_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_lug.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: lug
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_lug_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_orm.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: orm
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_orm_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_sot.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: sot
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_sot_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_swa.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: swa
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_swa_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_vai.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: vai
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_vai_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_wol.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: wol
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_wol_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_yaml ADDED
@@ -0,0 +1,33 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ tag:
2
+ - afrimgsm_tasks
3
+ - afrimgsm_tasks_prompt_4
4
+ dataset_path: masakhane/afrimgsm
5
+ output_type: generate_until
6
+ test_split: test
7
+ doc_to_target: '{% if answer is not none %}{{answer[21:]}}{% else %}{{answer_number|string}}{% endif %}'
8
+ target_delimiter: ""
9
+ generation_kwargs:
10
+ do_sample: false
11
+ until:
12
+ - 'Question:'
13
+ - </s>
14
+ - <|im_end|>
15
+ filter_list:
16
+ - name: remove_whitespace
17
+ filter:
18
+ - function: remove_whitespace
19
+ - function: take_first
20
+ - filter:
21
+ - function: regex
22
+ group_select: -1
23
+ regex_pattern: (-?[$0-9.,]{2,})|(-?[0-9]+)
24
+ - function: take_first
25
+ name: flexible-extract
26
+ metric_list:
27
+ - metric: exact_match
28
+ aggregation: mean
29
+ higher_is_better: true
30
+ ignore_case: true
31
+ ignore_punctuation: true
32
+ metadata:
33
+ version: 2.0
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_4/afrimgsm_yor.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: yor
3
+ doc_to_text: "Answer the given question with the appropriate numerical value, ensuring\
4
+ \ that the response is clear and without any supplementary information. \n\nQuestion:\
5
+ \ {{question}} \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_yor_prompt_4
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_amh.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: amh
3
+ doc_to_text: "For mathematical questions provided in Amharic language. Supply the\
4
+ \ accurate numeric answer to the provided question. \n\nQuestion: {{question}} \n\
5
+ Answer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_amh_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_eng.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: eng
3
+ doc_to_text: "For mathematical questions provided in English language. Supply the\
4
+ \ accurate numeric answer to the provided question. \n\nQuestion: {{question}} \n\
5
+ Answer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_eng_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_ewe.yaml ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: ewe
3
+ doc_to_text: "For mathematical questions provided in Ewe language. Supply the accurate\
4
+ \ numeric answer to the provided question. \n\nQuestion: {{question}} \nAnswer: "
5
+ include: afrimgsm_yaml
6
+ task: afrimgsm_ewe_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_fra.yaml ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: fra
3
+ doc_to_text: "For mathematical questions provided in French language. Supply the accurate\
4
+ \ numeric answer to the provided question. \n\nQuestion: {{question}} \nAnswer: "
5
+ include: afrimgsm_yaml
6
+ task: afrimgsm_fra_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_hau.yaml ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: hau
3
+ doc_to_text: "For mathematical questions provided in Hausa language. Supply the accurate\
4
+ \ numeric answer to the provided question. \n\nQuestion: {{question}} \nAnswer: "
5
+ include: afrimgsm_yaml
6
+ task: afrimgsm_hau_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_ibo.yaml ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: ibo
3
+ doc_to_text: "For mathematical questions provided in Igbo language. Supply the accurate\
4
+ \ numeric answer to the provided question. \n\nQuestion: {{question}} \nAnswer: "
5
+ include: afrimgsm_yaml
6
+ task: afrimgsm_ibo_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_kin.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: kin
3
+ doc_to_text: "For mathematical questions provided in Kinyarwanda language. Supply\
4
+ \ the accurate numeric answer to the provided question. \n\nQuestion: {{question}}\
5
+ \ \nAnswer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_kin_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_orm.yaml ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: orm
3
+ doc_to_text: "For mathematical questions provided in Oromo language. Supply the accurate\
4
+ \ numeric answer to the provided question. \n\nQuestion: {{question}} \nAnswer: "
5
+ include: afrimgsm_yaml
6
+ task: afrimgsm_orm_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_sna.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: sna
3
+ doc_to_text: "For mathematical questions provided in chiShona language. Supply the\
4
+ \ accurate numeric answer to the provided question. \n\nQuestion: {{question}} \n\
5
+ Answer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_sna_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_twi.yaml ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: twi
3
+ doc_to_text: "For mathematical questions provided in Twi language. Supply the accurate\
4
+ \ numeric answer to the provided question. \n\nQuestion: {{question}} \nAnswer: "
5
+ include: afrimgsm_yaml
6
+ task: afrimgsm_twi_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_wol.yaml ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: wol
3
+ doc_to_text: "For mathematical questions provided in Wolof language. Supply the accurate\
4
+ \ numeric answer to the provided question. \n\nQuestion: {{question}} \nAnswer: "
5
+ include: afrimgsm_yaml
6
+ task: afrimgsm_wol_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_xho.yaml ADDED
@@ -0,0 +1,7 @@
 
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: xho
3
+ doc_to_text: "For mathematical questions provided in isiXhosa language. Supply the\
4
+ \ accurate numeric answer to the provided question. \n\nQuestion: {{question}} \n\
5
+ Answer: "
6
+ include: afrimgsm_yaml
7
+ task: afrimgsm_xho_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_yaml ADDED
@@ -0,0 +1,33 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ tag:
2
+ - afrimgsm_tasks
3
+ - afrimgsm_tasks_prompt_5
4
+ dataset_path: masakhane/afrimgsm
5
+ output_type: generate_until
6
+ test_split: test
7
+ doc_to_target: '{% if answer is not none %}{{answer[21:]}}{% else %}{{answer_number|string}}{% endif %}'
8
+ target_delimiter: ""
9
+ generation_kwargs:
10
+ do_sample: false
11
+ until:
12
+ - 'Question:'
13
+ - </s>
14
+ - <|im_end|>
15
+ filter_list:
16
+ - name: remove_whitespace
17
+ filter:
18
+ - function: remove_whitespace
19
+ - function: take_first
20
+ - filter:
21
+ - function: regex
22
+ group_select: -1
23
+ regex_pattern: (-?[$0-9.,]{2,})|(-?[0-9]+)
24
+ - function: take_first
25
+ name: flexible-extract
26
+ metric_list:
27
+ - metric: exact_match
28
+ aggregation: mean
29
+ higher_is_better: true
30
+ ignore_case: true
31
+ ignore_punctuation: true
32
+ metadata:
33
+ version: 2.0
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct/prompt_5/afrimgsm_zul.yaml ADDED
@@ -0,0 +1,6 @@
 
 
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: zul
3
+ doc_to_text: "For mathematical questions provided in Zulu language. Supply the accurate\
4
+ \ numeric answer to the provided question. \n\nQuestion: {{question}} \nAnswer: "
5
+ include: afrimgsm_yaml
6
+ task: afrimgsm_zul_prompt_5
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/afrimgsm_cot.yaml ADDED
@@ -0,0 +1,9 @@
 
 
 
 
 
 
 
 
 
 
1
+ group: afrimgsm_cot-irokobench
2
+ task:
3
+ - afrimgsm_cot_tasks
4
+ aggregate_metric_list:
5
+ - metric: acc
6
+ aggregation: mean
7
+ weight_by_size: true
8
+ metadata:
9
+ version: 2
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_amh.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: amh
3
+ include: afrimgsm_cot_yaml
4
+ task: afrimgsm_cot_amh_prompt_1
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_fra.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: fra
3
+ include: afrimgsm_cot_yaml
4
+ task: afrimgsm_cot_fra_prompt_1
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_kin.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: kin
3
+ include: afrimgsm_cot_yaml
4
+ task: afrimgsm_cot_kin_prompt_1
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_lug.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: lug
3
+ include: afrimgsm_cot_yaml
4
+ task: afrimgsm_cot_lug_prompt_1
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_orm.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: orm
3
+ include: afrimgsm_cot_yaml
4
+ task: afrimgsm_cot_orm_prompt_1
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_sna.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: sna
3
+ include: afrimgsm_cot_yaml
4
+ task: afrimgsm_cot_sna_prompt_1
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_sot.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: sot
3
+ include: afrimgsm_cot_yaml
4
+ task: afrimgsm_cot_sot_prompt_1
lm-evaluation-harness/lm_eval/tasks/afrimgsm/direct_cot/prompt_1/afrimgsm_cot_swa.yaml ADDED
@@ -0,0 +1,4 @@
 
 
 
 
 
1
+ # Generated by utils.py
2
+ dataset_name: swa
3
+ include: afrimgsm_cot_yaml
4
+ task: afrimgsm_cot_swa_prompt_1