{ "_name_or_path": "/content/whisper_ur_csnorth_medium_model_v3/home/jovyan/upload-in-this/CS-training/CS-North-STT-training_data/Whisper_ur_csnorth_medium_model_v3", "activation_dropout": 0.0, "activation_function": "gelu", "apply_spec_augment": false, "architectures": [ "WhisperForConditionalGeneration" ], "attention_dropout": 0.0, "begin_suppress_tokens": null, "bos_token_id": 50257, "classifier_proj_size": 256, "d_model": 1024, "decoder_attention_heads": 16, "decoder_ffn_dim": 4096, "decoder_layerdrop": 0.0, "decoder_layers": 24, "decoder_start_token_id": 50258, "dropout": 0.0, "encoder_attention_heads": 16, "encoder_ffn_dim": 4096, "encoder_layerdrop": 0.0, "encoder_layers": 24, "eos_token_id": 50257, "forced_decoder_ids": null, "init_std": 0.02, "is_encoder_decoder": true, "mask_feature_length": 10, "mask_feature_min_masks": 0, "mask_feature_prob": 0.0, "mask_time_length": 10, "mask_time_min_masks": 2, "mask_time_prob": 0.05, "max_length": null, "max_source_positions": 1500, "max_target_positions": 448, "median_filter_width": 7, "model_type": "whisper", "num_hidden_layers": 24, "num_mel_bins": 80, "pad_token_id": 50257, "quantization_config": { "config_groups": { "group_0": { "input_activations": null, "output_activations": null, "targets": [ "Linear" ], "weights": { "actorder": null, "block_structure": null, "dynamic": false, "group_size": 128, "num_bits": 4, "observer": "minmax", "observer_kwargs": {}, "strategy": "group", "symmetric": true, "type": "int" } } }, "format": "pack-quantized", "global_compression_ratio": 2.5068917561966515, "ignore": [], "kv_cache_scheme": null, "quant_method": "compressed-tensors", "quantization_status": "compressed", "sparsity_config": { "format": "dense", "global_sparsity": 0.19957089243181236, "ignore": [ "model.encoder.layers.0.self_attn.k_proj", "model.encoder.layers.0.self_attn.v_proj", "model.encoder.layers.0.self_attn.q_proj", "model.encoder.layers.0.self_attn.out_proj", "model.encoder.layers.0.fc1", "model.encoder.layers.0.fc2", "model.encoder.layers.1.self_attn.k_proj", "model.encoder.layers.1.self_attn.v_proj", "model.encoder.layers.1.self_attn.q_proj", "model.encoder.layers.1.self_attn.out_proj", "model.encoder.layers.1.fc1", "model.encoder.layers.1.fc2", "model.encoder.layers.2.self_attn.k_proj", "model.encoder.layers.2.self_attn.v_proj", "model.encoder.layers.2.self_attn.q_proj", "model.encoder.layers.2.self_attn.out_proj", "model.encoder.layers.2.fc1", "model.encoder.layers.3.self_attn.k_proj", "model.encoder.layers.3.self_attn.v_proj", "model.encoder.layers.3.self_attn.q_proj", "model.encoder.layers.3.self_attn.out_proj", "model.encoder.layers.3.fc1", "model.encoder.layers.4.self_attn.k_proj", "model.encoder.layers.4.self_attn.v_proj", "model.encoder.layers.4.self_attn.q_proj", "model.encoder.layers.4.self_attn.out_proj", "model.encoder.layers.4.fc1", "model.encoder.layers.4.fc2", "model.encoder.layers.5.self_attn.k_proj", "model.encoder.layers.5.self_attn.v_proj", "model.encoder.layers.5.self_attn.q_proj", "model.encoder.layers.5.self_attn.out_proj", "model.encoder.layers.5.fc1", "model.encoder.layers.5.fc2", "model.encoder.layers.6.self_attn.k_proj", "model.encoder.layers.6.self_attn.v_proj", "model.encoder.layers.6.self_attn.q_proj", "model.encoder.layers.6.self_attn.out_proj", "model.encoder.layers.6.fc1", "model.encoder.layers.6.fc2", "model.encoder.layers.7.self_attn.k_proj", "model.encoder.layers.7.self_attn.v_proj", "model.encoder.layers.7.self_attn.q_proj", "model.encoder.layers.7.self_attn.out_proj", "model.encoder.layers.7.fc1", "model.encoder.layers.7.fc2", "model.encoder.layers.8.self_attn.k_proj", "model.encoder.layers.8.self_attn.v_proj", "model.encoder.layers.8.self_attn.q_proj", "model.encoder.layers.8.self_attn.out_proj", "model.encoder.layers.8.fc1", "model.encoder.layers.8.fc2", "model.encoder.layers.9.self_attn.k_proj", "model.encoder.layers.9.self_attn.v_proj", "model.encoder.layers.9.self_attn.q_proj", "model.encoder.layers.9.self_attn.out_proj", "model.encoder.layers.9.fc1", "model.encoder.layers.9.fc2", "model.encoder.layers.10.self_attn.k_proj", "model.encoder.layers.10.self_attn.v_proj", "model.encoder.layers.10.self_attn.q_proj", "model.encoder.layers.10.self_attn.out_proj", "model.encoder.layers.10.fc1", "model.encoder.layers.10.fc2", "model.encoder.layers.11.self_attn.k_proj", "model.encoder.layers.11.self_attn.v_proj", "model.encoder.layers.11.self_attn.q_proj", "model.encoder.layers.11.self_attn.out_proj", "model.encoder.layers.11.fc1", "model.encoder.layers.11.fc2", "model.encoder.layers.12.self_attn.k_proj", "model.encoder.layers.12.self_attn.v_proj", "model.encoder.layers.12.self_attn.q_proj", "model.encoder.layers.12.self_attn.out_proj", "model.encoder.layers.12.fc1", "model.encoder.layers.12.fc2", "model.encoder.layers.13.self_attn.k_proj", "model.encoder.layers.13.self_attn.v_proj", "model.encoder.layers.13.self_attn.q_proj", "model.encoder.layers.13.self_attn.out_proj", "model.encoder.layers.13.fc1", "model.encoder.layers.13.fc2", "model.encoder.layers.14.self_attn.k_proj", "model.encoder.layers.14.self_attn.v_proj", "model.encoder.layers.14.self_attn.q_proj", "model.encoder.layers.14.self_attn.out_proj", "model.encoder.layers.14.fc1", "model.encoder.layers.14.fc2", "model.encoder.layers.15.self_attn.k_proj", "model.encoder.layers.15.self_attn.v_proj", "model.encoder.layers.15.self_attn.q_proj", "model.encoder.layers.15.self_attn.out_proj", "model.encoder.layers.15.fc1", "model.encoder.layers.15.fc2", "model.encoder.layers.16.self_attn.k_proj", "model.encoder.layers.16.self_attn.v_proj", "model.encoder.layers.16.self_attn.q_proj", "model.encoder.layers.16.self_attn.out_proj", "model.encoder.layers.16.fc1", "model.encoder.layers.16.fc2", "model.encoder.layers.17.self_attn.k_proj", "model.encoder.layers.17.self_attn.v_proj", "model.encoder.layers.17.self_attn.q_proj", "model.encoder.layers.17.self_attn.out_proj", "model.encoder.layers.17.fc1", "model.encoder.layers.17.fc2", "model.encoder.layers.18.self_attn.k_proj", "model.encoder.layers.18.self_attn.v_proj", "model.encoder.layers.18.self_attn.q_proj", "model.encoder.layers.18.self_attn.out_proj", "model.encoder.layers.18.fc1", "model.encoder.layers.18.fc2", "model.encoder.layers.19.self_attn.k_proj", "model.encoder.layers.19.self_attn.v_proj", "model.encoder.layers.19.self_attn.q_proj", "model.encoder.layers.19.self_attn.out_proj", "model.encoder.layers.19.fc1", "model.encoder.layers.19.fc2", "model.encoder.layers.20.self_attn.k_proj", "model.encoder.layers.20.self_attn.v_proj", "model.encoder.layers.20.self_attn.q_proj", "model.encoder.layers.20.self_attn.out_proj", "model.encoder.layers.20.fc1", "model.encoder.layers.20.fc2", "model.encoder.layers.21.self_attn.k_proj", "model.encoder.layers.21.self_attn.v_proj", "model.encoder.layers.21.self_attn.q_proj", "model.encoder.layers.21.self_attn.out_proj", "model.encoder.layers.21.fc1", "model.encoder.layers.21.fc2", "model.encoder.layers.22.self_attn.k_proj", "model.encoder.layers.22.self_attn.v_proj", "model.encoder.layers.22.self_attn.q_proj", "model.encoder.layers.22.self_attn.out_proj", "model.encoder.layers.22.fc1", "model.encoder.layers.22.fc2", "model.encoder.layers.23.self_attn.k_proj", "model.encoder.layers.23.self_attn.v_proj", "model.encoder.layers.23.self_attn.q_proj", "model.encoder.layers.23.self_attn.out_proj", "model.encoder.layers.23.fc1", "model.encoder.layers.23.fc2", "model.decoder.layers.0.self_attn.k_proj", "model.decoder.layers.0.self_attn.v_proj", "model.decoder.layers.0.self_attn.q_proj", "model.decoder.layers.0.self_attn.out_proj", "model.decoder.layers.0.encoder_attn.k_proj", "model.decoder.layers.0.encoder_attn.v_proj", "model.decoder.layers.0.encoder_attn.q_proj", "model.decoder.layers.0.encoder_attn.out_proj", "model.decoder.layers.0.fc1", "model.decoder.layers.0.fc2", "model.decoder.layers.1.self_attn.k_proj", "model.decoder.layers.1.self_attn.v_proj", "model.decoder.layers.1.self_attn.q_proj", "model.decoder.layers.1.self_attn.out_proj", "model.decoder.layers.1.encoder_attn.k_proj", "model.decoder.layers.1.encoder_attn.v_proj", "model.decoder.layers.1.encoder_attn.q_proj", "model.decoder.layers.1.encoder_attn.out_proj", "model.decoder.layers.1.fc1", "model.decoder.layers.1.fc2", "model.decoder.layers.2.self_attn.k_proj", "model.decoder.layers.2.self_attn.v_proj", "model.decoder.layers.2.self_attn.q_proj", "model.decoder.layers.2.self_attn.out_proj", "model.decoder.layers.2.encoder_attn.k_proj", "model.decoder.layers.2.encoder_attn.v_proj", "model.decoder.layers.2.encoder_attn.q_proj", "model.decoder.layers.2.encoder_attn.out_proj", "model.decoder.layers.2.fc1", "model.decoder.layers.2.fc2", "model.decoder.layers.3.self_attn.k_proj", "model.decoder.layers.3.self_attn.v_proj", "model.decoder.layers.3.self_attn.q_proj", "model.decoder.layers.3.self_attn.out_proj", "model.decoder.layers.3.encoder_attn.k_proj", "model.decoder.layers.3.encoder_attn.v_proj", "model.decoder.layers.3.encoder_attn.q_proj", "model.decoder.layers.3.encoder_attn.out_proj", "model.decoder.layers.3.fc1", "model.decoder.layers.3.fc2", "model.decoder.layers.4.self_attn.k_proj", "model.decoder.layers.4.self_attn.v_proj", "model.decoder.layers.4.self_attn.q_proj", "model.decoder.layers.4.self_attn.out_proj", "model.decoder.layers.4.encoder_attn.k_proj", "model.decoder.layers.4.encoder_attn.v_proj", "model.decoder.layers.4.encoder_attn.q_proj", "model.decoder.layers.4.encoder_attn.out_proj", "model.decoder.layers.4.fc1", "model.decoder.layers.4.fc2", "model.decoder.layers.5.self_attn.k_proj", "model.decoder.layers.5.self_attn.v_proj", "model.decoder.layers.5.self_attn.q_proj", "model.decoder.layers.5.self_attn.out_proj", "model.decoder.layers.5.encoder_attn.k_proj", "model.decoder.layers.5.encoder_attn.v_proj", "model.decoder.layers.5.encoder_attn.q_proj", "model.decoder.layers.5.encoder_attn.out_proj", "model.decoder.layers.5.fc1", "model.decoder.layers.5.fc2", "model.decoder.layers.6.self_attn.k_proj", "model.decoder.layers.6.self_attn.v_proj", "model.decoder.layers.6.self_attn.q_proj", "model.decoder.layers.6.self_attn.out_proj", "model.decoder.layers.6.encoder_attn.k_proj", "model.decoder.layers.6.encoder_attn.v_proj", "model.decoder.layers.6.encoder_attn.q_proj", "model.decoder.layers.6.encoder_attn.out_proj", "model.decoder.layers.6.fc1", "model.decoder.layers.6.fc2", "model.decoder.layers.7.self_attn.k_proj", "model.decoder.layers.7.self_attn.v_proj", "model.decoder.layers.7.self_attn.q_proj", "model.decoder.layers.7.self_attn.out_proj", "model.decoder.layers.7.encoder_attn.k_proj", "model.decoder.layers.7.encoder_attn.v_proj", "model.decoder.layers.7.encoder_attn.q_proj", "model.decoder.layers.7.encoder_attn.out_proj", "model.decoder.layers.7.fc1", "model.decoder.layers.7.fc2", "model.decoder.layers.8.self_attn.k_proj", "model.decoder.layers.8.self_attn.v_proj", "model.decoder.layers.8.self_attn.q_proj", "model.decoder.layers.8.self_attn.out_proj", "model.decoder.layers.8.encoder_attn.k_proj", "model.decoder.layers.8.encoder_attn.v_proj", "model.decoder.layers.8.encoder_attn.q_proj", "model.decoder.layers.8.encoder_attn.out_proj", "model.decoder.layers.8.fc1", "model.decoder.layers.8.fc2", "model.decoder.layers.9.self_attn.k_proj", "model.decoder.layers.9.self_attn.v_proj", "model.decoder.layers.9.self_attn.q_proj", "model.decoder.layers.9.self_attn.out_proj", "model.decoder.layers.9.encoder_attn.k_proj", "model.decoder.layers.9.encoder_attn.v_proj", "model.decoder.layers.9.encoder_attn.q_proj", "model.decoder.layers.9.encoder_attn.out_proj", "model.decoder.layers.9.fc1", "model.decoder.layers.9.fc2", "model.decoder.layers.10.self_attn.k_proj", "model.decoder.layers.10.self_attn.v_proj", "model.decoder.layers.10.self_attn.q_proj", "model.decoder.layers.10.self_attn.out_proj", "model.decoder.layers.10.encoder_attn.k_proj", "model.decoder.layers.10.encoder_attn.v_proj", "model.decoder.layers.10.encoder_attn.q_proj", "model.decoder.layers.10.encoder_attn.out_proj", "model.decoder.layers.10.fc1", "model.decoder.layers.10.fc2", "model.decoder.layers.11.self_attn.k_proj", "model.decoder.layers.11.self_attn.v_proj", "model.decoder.layers.11.self_attn.q_proj", "model.decoder.layers.11.self_attn.out_proj", "model.decoder.layers.11.encoder_attn.k_proj", "model.decoder.layers.11.encoder_attn.v_proj", "model.decoder.layers.11.encoder_attn.q_proj", "model.decoder.layers.11.encoder_attn.out_proj", "model.decoder.layers.11.fc1", "model.decoder.layers.11.fc2", "model.decoder.layers.12.self_attn.k_proj", "model.decoder.layers.12.self_attn.v_proj", "model.decoder.layers.12.self_attn.q_proj", "model.decoder.layers.12.self_attn.out_proj", "model.decoder.layers.12.encoder_attn.k_proj", "model.decoder.layers.12.encoder_attn.v_proj", "model.decoder.layers.12.encoder_attn.q_proj", "model.decoder.layers.12.encoder_attn.out_proj", "model.decoder.layers.12.fc1", "model.decoder.layers.12.fc2", "model.decoder.layers.13.self_attn.k_proj", "model.decoder.layers.13.self_attn.v_proj", "model.decoder.layers.13.self_attn.q_proj", "model.decoder.layers.13.self_attn.out_proj", "model.decoder.layers.13.encoder_attn.k_proj", "model.decoder.layers.13.encoder_attn.v_proj", "model.decoder.layers.13.encoder_attn.q_proj", "model.decoder.layers.13.encoder_attn.out_proj", "model.decoder.layers.13.fc1", "model.decoder.layers.13.fc2", "model.decoder.layers.14.self_attn.k_proj", "model.decoder.layers.14.self_attn.v_proj", "model.decoder.layers.14.self_attn.q_proj", "model.decoder.layers.14.self_attn.out_proj", "model.decoder.layers.14.encoder_attn.k_proj", "model.decoder.layers.14.encoder_attn.v_proj", "model.decoder.layers.14.encoder_attn.q_proj", "model.decoder.layers.14.encoder_attn.out_proj", "model.decoder.layers.14.fc1", "model.decoder.layers.14.fc2", "model.decoder.layers.15.self_attn.k_proj", "model.decoder.layers.15.self_attn.v_proj", "model.decoder.layers.15.self_attn.q_proj", "model.decoder.layers.15.self_attn.out_proj", "model.decoder.layers.15.encoder_attn.k_proj", "model.decoder.layers.15.encoder_attn.v_proj", "model.decoder.layers.15.encoder_attn.q_proj", "model.decoder.layers.15.encoder_attn.out_proj", "model.decoder.layers.15.fc1", "model.decoder.layers.15.fc2", "model.decoder.layers.16.self_attn.k_proj", "model.decoder.layers.16.self_attn.v_proj", "model.decoder.layers.16.self_attn.q_proj", "model.decoder.layers.16.self_attn.out_proj", "model.decoder.layers.16.encoder_attn.k_proj", "model.decoder.layers.16.encoder_attn.v_proj", "model.decoder.layers.16.encoder_attn.q_proj", "model.decoder.layers.16.encoder_attn.out_proj", "model.decoder.layers.16.fc1", "model.decoder.layers.16.fc2", "model.decoder.layers.17.self_attn.k_proj", "model.decoder.layers.17.self_attn.v_proj", "model.decoder.layers.17.self_attn.q_proj", "model.decoder.layers.17.self_attn.out_proj", "model.decoder.layers.17.encoder_attn.k_proj", "model.decoder.layers.17.encoder_attn.v_proj", "model.decoder.layers.17.encoder_attn.q_proj", "model.decoder.layers.17.encoder_attn.out_proj", "model.decoder.layers.17.fc1", "model.decoder.layers.17.fc2", "model.decoder.layers.18.self_attn.k_proj", "model.decoder.layers.18.self_attn.v_proj", "model.decoder.layers.18.self_attn.q_proj", "model.decoder.layers.18.self_attn.out_proj", "model.decoder.layers.18.encoder_attn.k_proj", "model.decoder.layers.18.encoder_attn.v_proj", "model.decoder.layers.18.encoder_attn.q_proj", "model.decoder.layers.18.encoder_attn.out_proj", "model.decoder.layers.18.fc1", "model.decoder.layers.18.fc2", "model.decoder.layers.19.self_attn.k_proj", "model.decoder.layers.19.self_attn.v_proj", "model.decoder.layers.19.self_attn.q_proj", "model.decoder.layers.19.self_attn.out_proj", "model.decoder.layers.19.encoder_attn.k_proj", "model.decoder.layers.19.encoder_attn.v_proj", "model.decoder.layers.19.encoder_attn.q_proj", "model.decoder.layers.19.encoder_attn.out_proj", "model.decoder.layers.19.fc1", "model.decoder.layers.19.fc2", "model.decoder.layers.20.self_attn.k_proj", "model.decoder.layers.20.self_attn.v_proj", "model.decoder.layers.20.self_attn.q_proj", "model.decoder.layers.20.self_attn.out_proj", "model.decoder.layers.20.encoder_attn.k_proj", "model.decoder.layers.20.encoder_attn.v_proj", "model.decoder.layers.20.encoder_attn.q_proj", "model.decoder.layers.20.encoder_attn.out_proj", "model.decoder.layers.20.fc1", "model.decoder.layers.20.fc2", "model.decoder.layers.21.self_attn.k_proj", "model.decoder.layers.21.self_attn.v_proj", "model.decoder.layers.21.self_attn.q_proj", "model.decoder.layers.21.self_attn.out_proj", "model.decoder.layers.21.encoder_attn.k_proj", "model.decoder.layers.21.encoder_attn.v_proj", "model.decoder.layers.21.encoder_attn.q_proj", "model.decoder.layers.21.encoder_attn.out_proj", "model.decoder.layers.21.fc1", "model.decoder.layers.21.fc2", "model.decoder.layers.22.self_attn.k_proj", "model.decoder.layers.22.self_attn.v_proj", "model.decoder.layers.22.self_attn.q_proj", "model.decoder.layers.22.self_attn.out_proj", "model.decoder.layers.22.encoder_attn.k_proj", "model.decoder.layers.22.encoder_attn.v_proj", "model.decoder.layers.22.encoder_attn.q_proj", "model.decoder.layers.22.encoder_attn.out_proj", "model.decoder.layers.22.fc1", "model.decoder.layers.22.fc2", "model.decoder.layers.23.self_attn.k_proj", "model.decoder.layers.23.self_attn.v_proj", "model.decoder.layers.23.self_attn.q_proj", "model.decoder.layers.23.self_attn.out_proj", "model.decoder.layers.23.encoder_attn.k_proj", "model.decoder.layers.23.encoder_attn.v_proj", "model.decoder.layers.23.encoder_attn.q_proj", "model.decoder.layers.23.encoder_attn.out_proj", "model.decoder.layers.23.fc1", "model.decoder.layers.23.fc2", "proj_out" ], "registry_requires_subclass": false, "sparsity_structure": "unstructured", "targets": [ "model.encoder.layers.2.fc2", "model.encoder.layers.3.fc2" ] } }, "scale_embedding": false, "torch_dtype": "float32", "transformers_version": "4.48.3", "use_cache": true, "use_weighted_layer_sum": false, "vocab_size": 51865 }