Training in progress, epoch 1
Browse files- .gitattributes +2 -0
- .gitignore +1 -0
- config.json +46 -0
- pytorch_model.bin +3 -0
- runs/Aug25_08-31-38_8c03e379ea8d/1661416777.8043325/events.out.tfevents.1661416777.8c03e379ea8d.80.8 +3 -0
- runs/Aug25_08-41-27_8c03e379ea8d/1661416902.547067/events.out.tfevents.1661416902.8c03e379ea8d.912.1 +3 -0
- runs/Aug25_08-41-27_8c03e379ea8d/events.out.tfevents.1661416902.8c03e379ea8d.912.0 +3 -0
- runs/Aug25_08-42-02_8c03e379ea8d/1661416953.1945095/events.out.tfevents.1661416953.8c03e379ea8d.912.3 +3 -0
- runs/Aug25_08-42-02_8c03e379ea8d/events.out.tfevents.1661416953.8c03e379ea8d.912.2 +3 -0
- runs/Aug25_08-44-13_8c03e379ea8d/1661417067.0743074/events.out.tfevents.1661417067.8c03e379ea8d.1124.1 +3 -0
- runs/Aug25_08-44-13_8c03e379ea8d/events.out.tfevents.1661417067.8c03e379ea8d.1124.0 +3 -0
- source.spm +3 -0
- special_tokens_map.json +5 -0
- target.spm +3 -0
- tokenizer_config.json +13 -0
- training_args.bin +3 -0
- vocab.json +0 -0
.gitattributes
CHANGED
@@ -29,3 +29,5 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
|
|
29 |
*.zip filter=lfs diff=lfs merge=lfs -text
|
30 |
*.zst filter=lfs diff=lfs merge=lfs -text
|
31 |
*tfevents* filter=lfs diff=lfs merge=lfs -text
|
|
|
|
|
|
29 |
*.zip filter=lfs diff=lfs merge=lfs -text
|
30 |
*.zst filter=lfs diff=lfs merge=lfs -text
|
31 |
*tfevents* filter=lfs diff=lfs merge=lfs -text
|
32 |
+
source.spm filter=lfs diff=lfs merge=lfs -text
|
33 |
+
target.spm filter=lfs diff=lfs merge=lfs -text
|
.gitignore
ADDED
@@ -0,0 +1 @@
|
|
|
|
|
1 |
+
checkpoint-*/
|
config.json
ADDED
@@ -0,0 +1,46 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"_name_or_path": "Helsinki-NLP/opus-mt-tc-big-en-ko",
|
3 |
+
"activation_dropout": 0.0,
|
4 |
+
"activation_function": "relu",
|
5 |
+
"architectures": [
|
6 |
+
"MarianMTModel"
|
7 |
+
],
|
8 |
+
"attention_dropout": 0.0,
|
9 |
+
"bad_words_ids": [
|
10 |
+
[
|
11 |
+
32000
|
12 |
+
]
|
13 |
+
],
|
14 |
+
"bos_token_id": 0,
|
15 |
+
"classifier_dropout": 0.0,
|
16 |
+
"d_model": 1024,
|
17 |
+
"decoder_attention_heads": 16,
|
18 |
+
"decoder_ffn_dim": 4096,
|
19 |
+
"decoder_layerdrop": 0.0,
|
20 |
+
"decoder_layers": 6,
|
21 |
+
"decoder_start_token_id": 32000,
|
22 |
+
"decoder_vocab_size": 32001,
|
23 |
+
"dropout": 0.1,
|
24 |
+
"encoder_attention_heads": 16,
|
25 |
+
"encoder_ffn_dim": 4096,
|
26 |
+
"encoder_layerdrop": 0.0,
|
27 |
+
"encoder_layers": 6,
|
28 |
+
"eos_token_id": 2,
|
29 |
+
"forced_eos_token_id": 2,
|
30 |
+
"init_std": 0.02,
|
31 |
+
"is_encoder_decoder": true,
|
32 |
+
"max_length": 512,
|
33 |
+
"max_position_embeddings": 1024,
|
34 |
+
"model_type": "marian",
|
35 |
+
"normalize_embedding": false,
|
36 |
+
"num_beams": 4,
|
37 |
+
"num_hidden_layers": 6,
|
38 |
+
"pad_token_id": 32000,
|
39 |
+
"scale_embedding": true,
|
40 |
+
"share_encoder_decoder_embeddings": true,
|
41 |
+
"static_position_embeddings": true,
|
42 |
+
"torch_dtype": "float32",
|
43 |
+
"transformers_version": "4.21.2",
|
44 |
+
"use_cache": true,
|
45 |
+
"vocab_size": 32001
|
46 |
+
}
|
pytorch_model.bin
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:770d230a175b512d93ab26efa1f1d191880523ab293f898ce95d5b3cffe1a0db
|
3 |
+
size 836718085
|
runs/Aug25_08-31-38_8c03e379ea8d/1661416777.8043325/events.out.tfevents.1661416777.8c03e379ea8d.80.8
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:c60055f049b7a299a39edf9814567c43afa6b439b7d5a29584c645b0b0aaf2e6
|
3 |
+
size 5674
|
runs/Aug25_08-41-27_8c03e379ea8d/1661416902.547067/events.out.tfevents.1661416902.8c03e379ea8d.912.1
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:b31bb0315ce2899464804df6f3c0c410660b9028af2063cca155965839d81dfb
|
3 |
+
size 5674
|
runs/Aug25_08-41-27_8c03e379ea8d/events.out.tfevents.1661416902.8c03e379ea8d.912.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:a947fb6e00df85a576d6b493d6626f4a99b6b860046653ca35ba4d38d01deaaf
|
3 |
+
size 4285
|
runs/Aug25_08-42-02_8c03e379ea8d/1661416953.1945095/events.out.tfevents.1661416953.8c03e379ea8d.912.3
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:7b2740aab6475c6dfa773ad483506fa9936c283ad13d0c74244742f596d444e6
|
3 |
+
size 5674
|
runs/Aug25_08-42-02_8c03e379ea8d/events.out.tfevents.1661416953.8c03e379ea8d.912.2
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:006c0364f7933f4001a9453bfe2fb6afe5628747370629e8051c282f0ac56e19
|
3 |
+
size 4285
|
runs/Aug25_08-44-13_8c03e379ea8d/1661417067.0743074/events.out.tfevents.1661417067.8c03e379ea8d.1124.1
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:c2feeecc79846e4007cf858cbfb8d39d6d4181d919d8c9449f998b0b932c8de8
|
3 |
+
size 5674
|
runs/Aug25_08-44-13_8c03e379ea8d/events.out.tfevents.1661417067.8c03e379ea8d.1124.0
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:019977d44f0ef6a90b455ac1009ca5f4cc2999acee1c7d477baf99cce35da58d
|
3 |
+
size 5541
|
source.spm
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:3d0591e65c49541d82f48df33d7b322c3d4ee7aa0ee8747f9a7f9355dbf22c95
|
3 |
+
size 789870
|
special_tokens_map.json
ADDED
@@ -0,0 +1,5 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"eos_token": "</s>",
|
3 |
+
"pad_token": "<pad>",
|
4 |
+
"unk_token": "<unk>"
|
5 |
+
}
|
target.spm
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:3d2aa641a0890d8966ab8703b109895a4e522713ce99b4a0192bfacb495bc97c
|
3 |
+
size 815483
|
tokenizer_config.json
ADDED
@@ -0,0 +1,13 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1 |
+
{
|
2 |
+
"eos_token": "</s>",
|
3 |
+
"model_max_length": 512,
|
4 |
+
"name_or_path": "Helsinki-NLP/opus-mt-tc-big-en-ko",
|
5 |
+
"pad_token": "<pad>",
|
6 |
+
"separate_vocabs": false,
|
7 |
+
"source_lang": "en",
|
8 |
+
"sp_model_kwargs": {},
|
9 |
+
"special_tokens_map_file": null,
|
10 |
+
"target_lang": "ko",
|
11 |
+
"tokenizer_class": "MarianTokenizer",
|
12 |
+
"unk_token": "<unk>"
|
13 |
+
}
|
training_args.bin
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:f7a7a9090a75d1cbd143988e8cee2ceac6f95a5511f0eea38c7db42273f8bd47
|
3 |
+
size 3503
|
vocab.json
ADDED
The diff for this file is too large to render.
See raw diff
|
|