Upload folder using huggingface_hub

Files changed (3) hide show

config.json ADDED Viewed

+{
+  "architectures": [
+    "Gemma3nAudioEncoder"
+  ],
+  "conf_attention_chunk_size": 12,
+  "conf_attention_context_left": 13,
+  "conf_attention_context_right": 0,
+  "conf_attention_logit_cap": 50.0,
+  "conf_conv_kernel_size": 5,
+  "conf_num_attention_heads": 8,
+  "conf_num_hidden_layers": 12,
+  "conf_reduction_factor": 4,
+  "conf_residual_weight": 0.5,
+  "gradient_clipping": 10000000000.0,
+  "hidden_size": 1536,
+  "input_feat_size": 128,
+  "model_type": "gemma3n_audio",
+  "rms_norm_eps": 1e-06,
+  "sscp_conv_channel_size": [
+    128,
+    32
+  ],
+  "sscp_conv_group_norm_eps": 0.001,
+  "sscp_conv_kernel_size": [
+    [
+      3,
+      3
+    ],
+    [
+      3,
+      3
+    ]
+  ],
+  "sscp_conv_stride_size": [
+    [
+      2,
+      2
+    ],
+    [
+      2,
+      2
+    ]
+  ],
+  "torch_dtype": "float32",
+  "transformers_version": "4.53.2",
+  "vocab_offset": 262272,
+  "vocab_size": 128
+}

model.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:c07b48dc2c7e5bd5a41f74551c337dcfb2e1eaf34d78ad2c52739e05832a4942
+size 2725426224

preprocessor_config.json ADDED Viewed

+{
+  "crop_size": null,
+  "data_format": "channels_first",
+  "default_to_square": false,
+  "device": null,
+  "disable_grouping": null,
+  "dither": 0.0,
+  "do_center_crop": null,
+  "do_convert_rgb": null,
+  "do_normalize": false,
+  "do_rescale": true,
+  "do_resize": true,
+  "feature_extractor_type": "Gemma3nAudioFeatureExtractor",
+  "feature_size": 128,
+  "fft_length": 1024,
+  "fft_overdrive": true,
+  "frame_length": 512,
+  "hop_length": 160,
+  "image_mean": [
+    0.5,
+    0.5,
+    0.5
+  ],
+  "image_processor_type": "SiglipImageProcessorFast",
+  "image_seq_length": 256,
+  "image_std": [
+    0.5,
+    0.5,
+    0.5
+  ],
+  "input_data_format": null,
+  "input_scale_factor": 1.0,
+  "max_frequency": 7600.0,
+  "mel_floor": 1e-05,
+  "min_frequency": 125.0,
+  "padding_side": "right",
+  "padding_value": 0.0,
+  "per_bin_mean": null,
+  "per_bin_stddev": null,
+  "preemphasis": 0.97,
+  "preemphasis_htk_flavor": true,
+  "processor_class": "Gemma3nProcessor",
+  "resample": 2,
+  "rescale_factor": 0.00392156862745098,
+  "return_attention_mask": true,
+  "return_tensors": null,
+  "sampling_rate": 16000,
+  "size": {
+    "height": 768,
+    "width": 768
+  }
+}