Training in progress, step 500

Browse files

Files changed (10) hide show

.gitignore +1 -0
config.json +34 -0
pytorch_model.bin +3 -0
runs/Apr21_16-11-02_Sagi/1682115066.9454205/events.out.tfevents.1682115066.Sagi.13085.1 +3 -0
runs/Apr21_16-11-02_Sagi/events.out.tfevents.1682115066.Sagi.13085.0 +3 -0
special_tokens_map.json +1 -0
spiece.model +3 -0
tokenizer.json +0 -0
tokenizer_config.json +1 -0
training_args.bin +3 -0

.gitignore ADDED Viewed

	@@ -0,0 +1 @@


1	+ checkpoint-*/

config.json ADDED Viewed

	@@ -0,0 +1,34 @@

+{
+ "_name_or_path": "google/bigbird-roberta-base",
+ "architectures": [
+ "BigBirdForMultipleChoice"
+ ],
+ "attention_probs_dropout_prob": 0.1,
+ "attention_type": "block_sparse",
+ "block_size": 64,
+ "bos_token_id": 1,
+ "classifier_dropout": null,
+ "eos_token_id": 2,
+ "gradient_checkpointing": false,
+ "hidden_act": "gelu_new",
+ "hidden_dropout_prob": 0.1,
+ "hidden_size": 768,
+ "initializer_range": 0.02,
+ "intermediate_size": 3072,
+ "layer_norm_eps": 1e-12,
+ "max_position_embeddings": 4096,
+ "model_type": "big_bird",
+ "num_attention_heads": 12,
+ "num_hidden_layers": 12,
+ "num_random_blocks": 3,
+ "pad_token_id": 0,
+ "position_embedding_type": "absolute",
+ "rescale_embeddings": false,
+ "sep_token_id": 66,
+ "torch_dtype": "float32",
+ "transformers_version": "4.16.2",
+ "type_vocab_size": 2,
+ "use_bias": true,
+ "use_cache": true,
+ "vocab_size": 50358
+}

pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:4658a32e1c300b057023bff139d473befab8cc2810c2905f5ac6cc25526c4748
+size 509992821

runs/Apr21_16-11-02_Sagi/1682115066.9454205/events.out.tfevents.1682115066.Sagi.13085.1 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:8d35d61eda7c8f7440994ea1e4877ea448d46c30ee851d246d2a3d00bf9d5f36
+size 4833

runs/Apr21_16-11-02_Sagi/events.out.tfevents.1682115066.Sagi.13085.0 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:f257a57efd129feff928f3afe94b8bdf31f05854655226cde4d8d959913dfea3
+size 4054

special_tokens_map.json ADDED Viewed

	@@ -0,0 +1 @@

+ {"bos_token": {"content": "</s>", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true}, "eos_token": {"content": "<s>", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true}, "unk_token": {"content": "<unk>", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true}, "sep_token": {"content": "[SEP]", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true}, "pad_token": {"content": "<pad>", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true}, "cls_token": {"content": "[CLS]", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true}, "mask_token": {"content": "[MASK]", "single_word": false, "lstrip": true, "rstrip": false, "normalized": true}}

spiece.model ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:fdc81e1fc9d42e0c08b86d5b280d05d7c5e9747c4231c648f2b56b8e1d893c82
+size 845731

tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

tokenizer_config.json ADDED Viewed

	@@ -0,0 +1 @@

+ {"bos_token": {"content": "</s>", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true, "__type": "AddedToken"}, "eos_token": {"content": "<s>", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true, "__type": "AddedToken"}, "unk_token": {"content": "<unk>", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true, "__type": "AddedToken"}, "sep_token": {"content": "[SEP]", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true, "__type": "AddedToken"}, "pad_token": {"content": "<pad>", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true, "__type": "AddedToken"}, "cls_token": {"content": "[CLS]", "single_word": false, "lstrip": false, "rstrip": false, "normalized": true, "__type": "AddedToken"}, "mask_token": {"content": "[MASK]", "single_word": false, "lstrip": true, "rstrip": false, "normalized": true, "__type": "AddedToken"}, "model_max_length": 4096, "name_or_path": "google/bigbird-roberta-base", "special_tokens_map_file": "/home/sagi/.cache/huggingface/transformers/400be7e354ea6eb77319bcc7fa34899ec9fa2e3aff0fa677f6eb7e45a01b1548.75b358ecb30fa6b001d9d87bfde336c02d9123e7a8f5b90cc890d0f6efc3d4a3", "sp_model_kwargs": {}, "tokenizer_class": "BigBirdTokenizer"}

training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:b96fde85f5390147b5ad2bcf7a9658625d7ba1d8f3d8a88a92277b9d34b555f6
+size 3067