domenicrosati
commited on
Commit
·
a36b5ae
1
Parent(s):
012982b
Training in progress, epoch 1
Browse files- config.json +6 -6
- pytorch_model.bin +2 -2
- runs/Apr21_21-05-23_c118cb4fd099/1650575251.6404903/events.out.tfevents.1650575251.c118cb4fd099.73.1 +3 -0
- runs/Apr21_21-05-23_c118cb4fd099/events.out.tfevents.1650575251.c118cb4fd099.73.0 +3 -0
- runs/Apr21_21-15-59_c118cb4fd099/1650575768.0879936/events.out.tfevents.1650575768.c118cb4fd099.451.1 +3 -0
- runs/Apr21_21-15-59_c118cb4fd099/events.out.tfevents.1650575768.c118cb4fd099.451.0 +3 -0
- runs/Apr21_21-24-08_c118cb4fd099/1650576256.7896867/events.out.tfevents.1650576256.c118cb4fd099.693.1 +3 -0
- runs/Apr21_21-24-08_c118cb4fd099/events.out.tfevents.1650576256.c118cb4fd099.693.0 +3 -0
- tokenizer.json +1 -1
- tokenizer_config.json +1 -1
- training_args.bin +1 -1
config.json
CHANGED
@@ -1,11 +1,11 @@
|
|
1 |
{
|
2 |
-
"_name_or_path": "t5-
|
3 |
"architectures": [
|
4 |
"T5ForConditionalGeneration"
|
5 |
],
|
6 |
-
"d_ff":
|
7 |
"d_kv": 64,
|
8 |
-
"d_model":
|
9 |
"decoder_start_token_id": 0,
|
10 |
"dropout_rate": 0.1,
|
11 |
"eos_token_id": 1,
|
@@ -15,9 +15,9 @@
|
|
15 |
"layer_norm_epsilon": 1e-06,
|
16 |
"model_type": "t5",
|
17 |
"n_positions": 512,
|
18 |
-
"num_decoder_layers":
|
19 |
-
"num_heads":
|
20 |
-
"num_layers":
|
21 |
"output_past": true,
|
22 |
"pad_token_id": 0,
|
23 |
"relative_attention_max_distance": 128,
|
|
|
1 |
{
|
2 |
+
"_name_or_path": "t5-base",
|
3 |
"architectures": [
|
4 |
"T5ForConditionalGeneration"
|
5 |
],
|
6 |
+
"d_ff": 3072,
|
7 |
"d_kv": 64,
|
8 |
+
"d_model": 768,
|
9 |
"decoder_start_token_id": 0,
|
10 |
"dropout_rate": 0.1,
|
11 |
"eos_token_id": 1,
|
|
|
15 |
"layer_norm_epsilon": 1e-06,
|
16 |
"model_type": "t5",
|
17 |
"n_positions": 512,
|
18 |
+
"num_decoder_layers": 12,
|
19 |
+
"num_heads": 12,
|
20 |
+
"num_layers": 12,
|
21 |
"output_past": true,
|
22 |
"pad_token_id": 0,
|
23 |
"relative_attention_max_distance": 128,
|
pytorch_model.bin
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:4e954586461f5402ca618c13de38f6f652f8c287e425b4734047f536dc2dc4d0
|
3 |
+
size 891700799
|
runs/Apr21_21-05-23_c118cb4fd099/1650575251.6404903/events.out.tfevents.1650575251.c118cb4fd099.73.1
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:2dd24a6c107c7e577ce1a0e8b2561c94d42c13f1e5fb45e24ad5a044fdf7beb9
|
3 |
+
size 5087
|
runs/Apr21_21-05-23_c118cb4fd099/events.out.tfevents.1650575251.c118cb4fd099.73.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:d4bd00efbcdc0f831cd427ebd19407996d0f4ac4acff829d317c77b376035fe3
|
3 |
+
size 4229
|
runs/Apr21_21-15-59_c118cb4fd099/1650575768.0879936/events.out.tfevents.1650575768.c118cb4fd099.451.1
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:aff32ceb17d3b501439eae9cf4f7a47a2a97a7db87892d4dd066a68bcb36493c
|
3 |
+
size 5087
|
runs/Apr21_21-15-59_c118cb4fd099/events.out.tfevents.1650575768.c118cb4fd099.451.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:443ea6a68c948c072d3968858af5a3d738bf5be6c2a3c2cd70f6c9b70cf5f3fb
|
3 |
+
size 4229
|
runs/Apr21_21-24-08_c118cb4fd099/1650576256.7896867/events.out.tfevents.1650576256.c118cb4fd099.693.1
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:cfc32852ebc7eeb955902e8b9749aac2068a271a90afb3ed69ff10b928850826
|
3 |
+
size 5087
|
runs/Apr21_21-24-08_c118cb4fd099/events.out.tfevents.1650576256.c118cb4fd099.693.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:689cd15b0ac1737c0735493cea517c5609960b1d16a2dcdb9675a8e1330cf759
|
3 |
+
size 4699
|
tokenizer.json
CHANGED
@@ -2,7 +2,7 @@
|
|
2 |
"version": "1.0",
|
3 |
"truncation": {
|
4 |
"direction": "Right",
|
5 |
-
"max_length":
|
6 |
"strategy": "LongestFirst",
|
7 |
"stride": 0
|
8 |
},
|
|
|
2 |
"version": "1.0",
|
3 |
"truncation": {
|
4 |
"direction": "Right",
|
5 |
+
"max_length": 128,
|
6 |
"strategy": "LongestFirst",
|
7 |
"stride": 0
|
8 |
},
|
tokenizer_config.json
CHANGED
@@ -1 +1 @@
|
|
1 |
-
{"eos_token": "</s>", "unk_token": "<unk>", "pad_token": "<pad>", "extra_ids": 100, "additional_special_tokens": ["<extra_id_0>", "<extra_id_1>", "<extra_id_2>", "<extra_id_3>", "<extra_id_4>", "<extra_id_5>", "<extra_id_6>", "<extra_id_7>", "<extra_id_8>", "<extra_id_9>", "<extra_id_10>", "<extra_id_11>", "<extra_id_12>", "<extra_id_13>", "<extra_id_14>", "<extra_id_15>", "<extra_id_16>", "<extra_id_17>", "<extra_id_18>", "<extra_id_19>", "<extra_id_20>", "<extra_id_21>", "<extra_id_22>", "<extra_id_23>", "<extra_id_24>", "<extra_id_25>", "<extra_id_26>", "<extra_id_27>", "<extra_id_28>", "<extra_id_29>", "<extra_id_30>", "<extra_id_31>", "<extra_id_32>", "<extra_id_33>", "<extra_id_34>", "<extra_id_35>", "<extra_id_36>", "<extra_id_37>", "<extra_id_38>", "<extra_id_39>", "<extra_id_40>", "<extra_id_41>", "<extra_id_42>", "<extra_id_43>", "<extra_id_44>", "<extra_id_45>", "<extra_id_46>", "<extra_id_47>", "<extra_id_48>", "<extra_id_49>", "<extra_id_50>", "<extra_id_51>", "<extra_id_52>", "<extra_id_53>", "<extra_id_54>", "<extra_id_55>", "<extra_id_56>", "<extra_id_57>", "<extra_id_58>", "<extra_id_59>", "<extra_id_60>", "<extra_id_61>", "<extra_id_62>", "<extra_id_63>", "<extra_id_64>", "<extra_id_65>", "<extra_id_66>", "<extra_id_67>", "<extra_id_68>", "<extra_id_69>", "<extra_id_70>", "<extra_id_71>", "<extra_id_72>", "<extra_id_73>", "<extra_id_74>", "<extra_id_75>", "<extra_id_76>", "<extra_id_77>", "<extra_id_78>", "<extra_id_79>", "<extra_id_80>", "<extra_id_81>", "<extra_id_82>", "<extra_id_83>", "<extra_id_84>", "<extra_id_85>", "<extra_id_86>", "<extra_id_87>", "<extra_id_88>", "<extra_id_89>", "<extra_id_90>", "<extra_id_91>", "<extra_id_92>", "<extra_id_93>", "<extra_id_94>", "<extra_id_95>", "<extra_id_96>", "<extra_id_97>", "<extra_id_98>", "<extra_id_99>"], "return_tensors": "tf", "model_max_length": 512, "special_tokens_map_file": null, "name_or_path": "t5-
|
|
|
1 |
+
{"eos_token": "</s>", "unk_token": "<unk>", "pad_token": "<pad>", "extra_ids": 100, "additional_special_tokens": ["<extra_id_0>", "<extra_id_1>", "<extra_id_2>", "<extra_id_3>", "<extra_id_4>", "<extra_id_5>", "<extra_id_6>", "<extra_id_7>", "<extra_id_8>", "<extra_id_9>", "<extra_id_10>", "<extra_id_11>", "<extra_id_12>", "<extra_id_13>", "<extra_id_14>", "<extra_id_15>", "<extra_id_16>", "<extra_id_17>", "<extra_id_18>", "<extra_id_19>", "<extra_id_20>", "<extra_id_21>", "<extra_id_22>", "<extra_id_23>", "<extra_id_24>", "<extra_id_25>", "<extra_id_26>", "<extra_id_27>", "<extra_id_28>", "<extra_id_29>", "<extra_id_30>", "<extra_id_31>", "<extra_id_32>", "<extra_id_33>", "<extra_id_34>", "<extra_id_35>", "<extra_id_36>", "<extra_id_37>", "<extra_id_38>", "<extra_id_39>", "<extra_id_40>", "<extra_id_41>", "<extra_id_42>", "<extra_id_43>", "<extra_id_44>", "<extra_id_45>", "<extra_id_46>", "<extra_id_47>", "<extra_id_48>", "<extra_id_49>", "<extra_id_50>", "<extra_id_51>", "<extra_id_52>", "<extra_id_53>", "<extra_id_54>", "<extra_id_55>", "<extra_id_56>", "<extra_id_57>", "<extra_id_58>", "<extra_id_59>", "<extra_id_60>", "<extra_id_61>", "<extra_id_62>", "<extra_id_63>", "<extra_id_64>", "<extra_id_65>", "<extra_id_66>", "<extra_id_67>", "<extra_id_68>", "<extra_id_69>", "<extra_id_70>", "<extra_id_71>", "<extra_id_72>", "<extra_id_73>", "<extra_id_74>", "<extra_id_75>", "<extra_id_76>", "<extra_id_77>", "<extra_id_78>", "<extra_id_79>", "<extra_id_80>", "<extra_id_81>", "<extra_id_82>", "<extra_id_83>", "<extra_id_84>", "<extra_id_85>", "<extra_id_86>", "<extra_id_87>", "<extra_id_88>", "<extra_id_89>", "<extra_id_90>", "<extra_id_91>", "<extra_id_92>", "<extra_id_93>", "<extra_id_94>", "<extra_id_95>", "<extra_id_96>", "<extra_id_97>", "<extra_id_98>", "<extra_id_99>"], "return_tensors": "tf", "model_max_length": 512, "special_tokens_map_file": null, "name_or_path": "t5-base", "tokenizer_class": "T5Tokenizer"}
|
training_args.bin
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 3183
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:fa9dcb0f67c6c2fce02d52708fbc28a138510d2652206eed93745ecd02e9efa0
|
3 |
size 3183
|