model: add new checkpoint for GC4 ELECTRA model
Browse files- .gitattributes +3 -0
- config.json +27 -0
- model.ckpt-400000.data-00000-of-00001 +3 -0
- model.ckpt-400000.index +3 -0
- model.ckpt-400000.meta +3 -0
- pytorch_model.bin +3 -0
- tf_model.h5 +3 -0
    	
        .gitattributes
    CHANGED
    
    | @@ -14,3 +14,6 @@ | |
| 14 | 
             
            *.pb filter=lfs diff=lfs merge=lfs -text
         | 
| 15 | 
             
            *.pt filter=lfs diff=lfs merge=lfs -text
         | 
| 16 | 
             
            *.pth filter=lfs diff=lfs merge=lfs -text
         | 
|  | |
|  | |
|  | 
|  | |
| 14 | 
             
            *.pb filter=lfs diff=lfs merge=lfs -text
         | 
| 15 | 
             
            *.pt filter=lfs diff=lfs merge=lfs -text
         | 
| 16 | 
             
            *.pth filter=lfs diff=lfs merge=lfs -text
         | 
| 17 | 
            +
            *.meta filter=lfs diff=lfs merge=lfs -text
         | 
| 18 | 
            +
            *.data-00000-of-00001 filter=lfs diff=lfs merge=lfs -text
         | 
| 19 | 
            +
            *.index filter=lfs diff=lfs merge=lfs -text
         | 
    	
        config.json
    ADDED
    
    | @@ -0,0 +1,27 @@ | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            {
         | 
| 2 | 
            +
              "_name_or_path": "electra-base-gc4-64k-400000-cased-generator",
         | 
| 3 | 
            +
              "architectures": [
         | 
| 4 | 
            +
                "ElectraForMaskedLM"
         | 
| 5 | 
            +
              ],
         | 
| 6 | 
            +
              "attention_probs_dropout_prob": 0.1,
         | 
| 7 | 
            +
              "embedding_size": 768,
         | 
| 8 | 
            +
              "hidden_act": "gelu",
         | 
| 9 | 
            +
              "hidden_dropout_prob": 0.1,
         | 
| 10 | 
            +
              "hidden_size": 256,
         | 
| 11 | 
            +
              "initializer_range": 0.02,
         | 
| 12 | 
            +
              "intermediate_size": 1024,
         | 
| 13 | 
            +
              "layer_norm_eps": 1e-12,
         | 
| 14 | 
            +
              "max_position_embeddings": 512,
         | 
| 15 | 
            +
              "model_type": "electra",
         | 
| 16 | 
            +
              "num_attention_heads": 4,
         | 
| 17 | 
            +
              "num_hidden_layers": 12,
         | 
| 18 | 
            +
              "pad_token_id": 0,
         | 
| 19 | 
            +
              "position_embedding_type": "absolute",
         | 
| 20 | 
            +
              "summary_activation": "gelu",
         | 
| 21 | 
            +
              "summary_last_dropout": 0.1,
         | 
| 22 | 
            +
              "summary_type": "first",
         | 
| 23 | 
            +
              "summary_use_proj": true,
         | 
| 24 | 
            +
              "transformers_version": "4.6.0.dev0",
         | 
| 25 | 
            +
              "type_vocab_size": 2,
         | 
| 26 | 
            +
              "vocab_size": 64000
         | 
| 27 | 
            +
            }
         | 
    	
        model.ckpt-400000.data-00000-of-00001
    ADDED
    
    | @@ -0,0 +1,3 @@ | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            version https://git-lfs.github.com/spec/v1
         | 
| 2 | 
            +
            oid sha256:7714c4301fa75d26fc2d5a428d50143e8be98444de2a54712ca649f6d6b4ae17
         | 
| 3 | 
            +
            size 1741572116
         | 
    	
        model.ckpt-400000.index
    ADDED
    
    | @@ -0,0 +1,3 @@ | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            version https://git-lfs.github.com/spec/v1
         | 
| 2 | 
            +
            oid sha256:a399ad0ae6779c35e002c4371f82b8cbb469d4c0f0ff29c41366c2a8b50b40fa
         | 
| 3 | 
            +
            size 17988
         | 
    	
        model.ckpt-400000.meta
    ADDED
    
    | @@ -0,0 +1,3 @@ | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            version https://git-lfs.github.com/spec/v1
         | 
| 2 | 
            +
            oid sha256:3957e25f9af858416ca5eb550cc852f4eba89d46ec8089efcbf57daeaa458238
         | 
| 3 | 
            +
            size 9227280
         | 
    	
        pytorch_model.bin
    ADDED
    
    | @@ -0,0 +1,3 @@ | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            version https://git-lfs.github.com/spec/v1
         | 
| 2 | 
            +
            oid sha256:f24d93189297d454a1eac80f2996b7fc92412d8053fba1415da3216eb40a10b3
         | 
| 3 | 
            +
            size 238029732
         | 
    	
        tf_model.h5
    ADDED
    
    | @@ -0,0 +1,3 @@ | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            version https://git-lfs.github.com/spec/v1
         | 
| 2 | 
            +
            oid sha256:5180889d8dd15f57e51095f1baa9c895bf1f20247fd71bb1fd429409b5af98ce
         | 
| 3 | 
            +
            size 436414984
         | 
