Training in progress, step 500
Browse files
    	
        adapter_config.json
    ADDED
    
    | @@ -0,0 +1,28 @@ | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            {
         | 
| 2 | 
            +
              "alpha_pattern": {},
         | 
| 3 | 
            +
              "auto_mapping": null,
         | 
| 4 | 
            +
              "base_model_name_or_path": "bertin-project/BOLETIN",
         | 
| 5 | 
            +
              "bias": "none",
         | 
| 6 | 
            +
              "fan_in_fan_out": false,
         | 
| 7 | 
            +
              "inference_mode": true,
         | 
| 8 | 
            +
              "init_lora_weights": true,
         | 
| 9 | 
            +
              "layers_pattern": null,
         | 
| 10 | 
            +
              "layers_to_transform": null,
         | 
| 11 | 
            +
              "loftq_config": {},
         | 
| 12 | 
            +
              "lora_alpha": 32,
         | 
| 13 | 
            +
              "lora_dropout": 0.1,
         | 
| 14 | 
            +
              "megatron_config": null,
         | 
| 15 | 
            +
              "megatron_core": "megatron.core",
         | 
| 16 | 
            +
              "modules_to_save": null,
         | 
| 17 | 
            +
              "peft_type": "LORA",
         | 
| 18 | 
            +
              "r": 16,
         | 
| 19 | 
            +
              "rank_pattern": {},
         | 
| 20 | 
            +
              "revision": null,
         | 
| 21 | 
            +
              "target_modules": [
         | 
| 22 | 
            +
                "k_proj",
         | 
| 23 | 
            +
                "o_proj",
         | 
| 24 | 
            +
                "q_proj",
         | 
| 25 | 
            +
                "v_proj"
         | 
| 26 | 
            +
              ],
         | 
| 27 | 
            +
              "task_type": "CAUSAL_LM"
         | 
| 28 | 
            +
            }
         | 
    	
        adapter_model.safetensors
    ADDED
    
    | @@ -0,0 +1,3 @@ | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            version https://git-lfs.github.com/spec/v1
         | 
| 2 | 
            +
            oid sha256:f525e4b2e33fbe660f04adeb6b6d84b604c9a97487d88a68417ffd3b51cc838b
         | 
| 3 | 
            +
            size 44062096
         | 
    	
        runs/Feb09_10-53-59_dante/events.out.tfevents.1707472446.dante.1997872.0
    ADDED
    
    | @@ -0,0 +1,3 @@ | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            version https://git-lfs.github.com/spec/v1
         | 
| 2 | 
            +
            oid sha256:0bad15fdf887ecea51bfe7e34027e0b1da559f4ba38c458fd3666422106ee531
         | 
| 3 | 
            +
            size 83245
         | 
    	
        training_args.bin
    ADDED
    
    | @@ -0,0 +1,3 @@ | |
|  | |
|  | |
|  | 
|  | |
| 1 | 
            +
            version https://git-lfs.github.com/spec/v1
         | 
| 2 | 
            +
            oid sha256:5e69467ac6509fdb714bd0fc694b51ebafcfbf1dc624bedd93f485396743b272
         | 
| 3 | 
            +
            size 4728
         | 
