bowenbaoamd commited on Oct 18, 2024

Commit

22bcd0c

verified ·

1 Parent(s): 3e3f59e

Upload folder using huggingface_hub

Browse files

Files changed (26) hide show

config.json +230 -52
generation_config.json +2 -1
model-00001-of-00019.safetensors +2 -2
model-00002-of-00019.safetensors +2 -2
model-00003-of-00019.safetensors +2 -2
model-00004-of-00019.safetensors +2 -2
model-00005-of-00019.safetensors +2 -2
model-00006-of-00019.safetensors +2 -2
model-00007-of-00019.safetensors +2 -2
model-00008-of-00019.safetensors +2 -2
model-00009-of-00019.safetensors +2 -2
model-00010-of-00019.safetensors +2 -2
model-00011-of-00019.safetensors +2 -2
model-00012-of-00019.safetensors +2 -2
model-00013-of-00019.safetensors +2 -2
model-00014-of-00019.safetensors +2 -2
model-00015-of-00019.safetensors +2 -2
model-00016-of-00019.safetensors +2 -2
model-00017-of-00019.safetensors +2 -2
model-00018-of-00019.safetensors +2 -2
model-00019-of-00019.safetensors +2 -2
model.safetensors.index.json +0 -0
preprocessor_config.json +25 -0
special_tokens_map.json +1 -7
tokenizer.json +2 -2
tokenizer_config.json +3 -2

config.json CHANGED Viewed

@@ -1,46 +1,10 @@
 {
   "architectures": [
-    "MllamaForCausalLM"
   ],
-  "bos_token_id": 128000,
-  "cross_attention_layers": [
-    3,
-    8,
-    13,
-    18,
-    23,
-    28,
-    33,
-    38,
-    43,
-    48,
-    53,
-    58,
-    63,
-    68,
-    73,
-    78,
-    83,
-    88,
-    93,
-    98
-  ],
-  "dropout": 0,
-  "eos_token_id": [
-    128001,
-    128008,
-    128009
-  ],
-  "hidden_act": "silu",
-  "hidden_size": 8192,
-  "initializer_range": 0.02,
-  "intermediate_size": 28672,
-  "max_position_embeddings": 131072,
-  "model_type": "mllama_text_model",
-  "num_attention_heads": 64,
-  "num_hidden_layers": 100,
-  "num_key_value_heads": 8,
-  "pad_token_id": 128004,
   "quantization_config": {
     "activation_scheme": "static",
     "ignored_layers": [
@@ -49,18 +13,232 @@
     "kv_cache_scheme": "static",
     "quant_method": "fp8"
   },
-  "rms_norm_eps": 1e-05,
-  "rope_scaling": {
-    "factor": 8.0,
-    "high_freq_factor": 4.0,
-    "low_freq_factor": 1.0,
-    "original_max_position_embeddings": 8192,
-    "rope_type": "llama3"
   },
-  "rope_theta": 500000.0,
-  "tie_word_embeddings": false,
   "torch_dtype": "bfloat16",
-  "transformers_version": "4.45.1",
-  "use_cache": true,
-  "vocab_size": 128256
 }

 {
+  "_name_or_path": "/model_path/meta-llama/Llama-3.2-90B-Vision-Instruct/",
   "architectures": [
+    "MllamaForConditionalGeneration"
   ],
+  "image_token_index": 128256,
+  "model_type": "mllama",
   "quantization_config": {
     "activation_scheme": "static",
     "ignored_layers": [
     "kv_cache_scheme": "static",
     "quant_method": "fp8"
   },
+  "text_config": {
+    "_name_or_path": "",
+    "add_cross_attention": false,
+    "architectures": null,
+    "bad_words_ids": null,
+    "begin_suppress_tokens": null,
+    "bos_token_id": 128000,
+    "chunk_size_feed_forward": 0,
+    "cross_attention_hidden_size": null,
+    "cross_attention_layers": [
+      3,
+      8,
+      13,
+      18,
+      23,
+      28,
+      33,
+      38,
+      43,
+      48,
+      53,
+      58,
+      63,
+      68,
+      73,
+      78,
+      83,
+      88,
+      93,
+      98
+    ],
+    "decoder_start_token_id": null,
+    "diversity_penalty": 0.0,
+    "do_sample": false,
+    "dropout": 0,
+    "early_stopping": false,
+    "encoder_no_repeat_ngram_size": 0,
+    "eos_token_id": [
+      128001,
+      128008,
+      128009
+    ],
+    "exponential_decay_length_penalty": null,
+    "finetuning_task": null,
+    "forced_bos_token_id": null,
+    "forced_eos_token_id": null,
+    "hidden_act": "silu",
+    "hidden_size": 8192,
+    "id2label": {
+      "0": "LABEL_0",
+      "1": "LABEL_1"
+    },
+    "initializer_range": 0.02,
+    "intermediate_size": 28672,
+    "is_decoder": false,
+    "is_encoder_decoder": false,
+    "label2id": {
+      "LABEL_0": 0,
+      "LABEL_1": 1
+    },
+    "length_penalty": 1.0,
+    "max_length": 20,
+    "max_position_embeddings": 131072,
+    "min_length": 0,
+    "model_type": "mllama_text_model",
+    "no_repeat_ngram_size": 0,
+    "num_attention_heads": 64,
+    "num_beam_groups": 1,
+    "num_beams": 1,
+    "num_hidden_layers": 100,
+    "num_key_value_heads": 8,
+    "num_return_sequences": 1,
+    "output_attentions": false,
+    "output_hidden_states": false,
+    "output_scores": false,
+    "pad_token_id": 128004,
+    "prefix": null,
+    "problem_type": null,
+    "pruned_heads": {},
+    "remove_invalid_values": false,
+    "repetition_penalty": 1.0,
+    "return_dict": true,
+    "return_dict_in_generate": false,
+    "rms_norm_eps": 1e-05,
+    "rope_scaling": {
+      "factor": 8.0,
+      "high_freq_factor": 4.0,
+      "low_freq_factor": 1.0,
+      "original_max_position_embeddings": 8192,
+      "rope_type": "llama3"
+    },
+    "rope_theta": 500000.0,
+    "sep_token_id": null,
+    "suppress_tokens": null,
+    "task_specific_params": null,
+    "temperature": 1.0,
+    "tf_legacy_loss": false,
+    "tie_encoder_decoder": false,
+    "tie_word_embeddings": false,
+    "tokenizer_class": null,
+    "top_k": 50,
+    "top_p": 1.0,
+    "torch_dtype": "bfloat16",
+    "torchscript": false,
+    "typical_p": 1.0,
+    "use_bfloat16": false,
+    "use_cache": true,
+    "vocab_size": 128256
   },
   "torch_dtype": "bfloat16",
+  "transformers_version": "4.45.2",
+  "vision_config": {
+    "_name_or_path": "",
+    "add_cross_attention": false,
+    "architectures": null,
+    "attention_heads": 16,
+    "bad_words_ids": null,
+    "begin_suppress_tokens": null,
+    "bos_token_id": null,
+    "chunk_size_feed_forward": 0,
+    "cross_attention_hidden_size": null,
+    "decoder_start_token_id": null,
+    "diversity_penalty": 0.0,
+    "do_sample": false,
+    "early_stopping": false,
+    "encoder_no_repeat_ngram_size": 0,
+    "eos_token_id": null,
+    "exponential_decay_length_penalty": null,
+    "finetuning_task": null,
+    "forced_bos_token_id": null,
+    "forced_eos_token_id": null,
+    "hidden_act": "gelu",
+    "hidden_size": 1280,
+    "id2label": {
+      "0": "LABEL_0",
+      "1": "LABEL_1"
+    },
+    "image_size": 560,
+    "initializer_range": 0.02,
+    "intermediate_layers_indices": [
+      3,
+      7,
+      15,
+      23,
+      30
+    ],
+    "intermediate_size": 5120,
+    "is_decoder": false,
+    "is_encoder_decoder": false,
+    "label2id": {
+      "LABEL_0": 0,
+      "LABEL_1": 1
+    },
+    "length_penalty": 1.0,
+    "max_length": 20,
+    "max_num_tiles": 4,
+    "min_length": 0,
+    "model_type": "mllama_vision_model",
+    "no_repeat_ngram_size": 0,
+    "norm_eps": 1e-05,
+    "num_beam_groups": 1,
+    "num_beams": 1,
+    "num_channels": 3,
+    "num_global_layers": 8,
+    "num_hidden_layers": 32,
+    "num_return_sequences": 1,
+    "output_attentions": false,
+    "output_hidden_states": false,
+    "output_scores": false,
+    "pad_token_id": null,
+    "patch_size": 14,
+    "prefix": null,
+    "problem_type": null,
+    "pruned_heads": {},
+    "remove_invalid_values": false,
+    "repetition_penalty": 1.0,
+    "return_dict": true,
+    "return_dict_in_generate": false,
+    "sep_token_id": null,
+    "supported_aspect_ratios": [
+      [
+        1,
+        1
+      ],
+      [
+        1,
+        2
+      ],
+      [
+        1,
+        3
+      ],
+      [
+        1,
+        4
+      ],
+      [
+        2,
+        1
+      ],
+      [
+        2,
+        2
+      ],
+      [
+        3,
+        1
+      ],
+      [
+        4,
+        1
+      ]
+    ],
+    "suppress_tokens": null,
+    "task_specific_params": null,
+    "temperature": 1.0,
+    "tf_legacy_loss": false,
+    "tie_encoder_decoder": false,
+    "tie_word_embeddings": true,
+    "tokenizer_class": null,
+    "top_k": 50,
+    "top_p": 1.0,
+    "torch_dtype": "bfloat16",
+    "torchscript": false,
+    "typical_p": 1.0,
+    "use_bfloat16": false,
+    "vision_output_dim": 7680
+  }
 }

generation_config.json CHANGED Viewed

@@ -1,4 +1,5 @@
 {
   "bos_token_id": 128000,
   "do_sample": true,
   "eos_token_id": [
@@ -9,5 +10,5 @@
   "pad_token_id": 128004,
   "temperature": 0.6,
   "top_p": 0.9,
-  "transformers_version": "4.45.1"
 }

 {
+  "attn_implementation": "eager",
   "bos_token_id": 128000,
   "do_sample": true,
   "eos_token_id": [
   "pad_token_id": 128004,
   "temperature": 0.6,
   "top_p": 0.9,
+  "transformers_version": "4.45.2"
 }

model-00001-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:0bd0a1fcc14896bb9b5c630e874e2df046065835bcdfbe557d9d8dff5dd85768
-size 4819511324

 version https://git-lfs.github.com/spec/v1
+oid sha256:a5633ead602ca8cd3f5603a1b352f916379f37bf7c92a9454a73c91b3d425bc4
+size 4835349758

model-00002-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:182efcf174d3e7f0c0aa573c1aebefc2c63a2f98760d89592298c80477ab2930
-size 4983028204

 version https://git-lfs.github.com/spec/v1
+oid sha256:c3c42c5c10a552ab129ffa92fc6a5a2b07a639920d9497f9d87ed52452a90e9d
+size 4983046716

model-00003-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:e95e09a2ba51f63651d40488135675f49492593e0e940d87fd8dc1ca69d870c5
-size 4899126616

 version https://git-lfs.github.com/spec/v1
+oid sha256:8237d40f2800f714a402d70c2e4bcdeb783e774b0af83996e17abdf27adacd14
+size 4899128696

model-00004-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:0d7a829e6b0b4ee3351485349a605608a1dba6103f83150c63e1e833cdb9e166
-size 4899159732

 version https://git-lfs.github.com/spec/v1
+oid sha256:81b197e264111aa05d4ffc5653659cc7bd0d0e35cf42da153d388225f6709844
+size 4899179408

model-00005-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:6fada21ac6d0cc5850b8a2e4f567a7f0bd9786b383440204a97b8119c3b3e0bf
-size 4899159724

 version https://git-lfs.github.com/spec/v1
+oid sha256:48c7dfa55ab23d7b2137a76d9f869f72c2e57635eca9971099e2f32021ab86cb
+size 4899145404

model-00006-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:2aab0c845d1df5f5539cc24ab2bf3bd7123551a89daa3150a8780fbfb1df5772
-size 4983044812

 version https://git-lfs.github.com/spec/v1
+oid sha256:5ae76eab5c8ec24d102280131ca0e5b08dfc5b13a1b365dfeb794e018f4d241e
+size 4983046860

model-00007-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:c69289f2379fa24b4fb90ef41979f58eec7f59a7eba9110710c60be67fd058b3
-size 4899126632

 version https://git-lfs.github.com/spec/v1
+oid sha256:4b75365ead7409e54bcdedb9796a50fcb99b206073e7f5dede2ff34761c4ff05
+size 4899128768

model-00008-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:8eb8aaa066f144dbd2de612b93e0e3a217dcbca6f7e8d3181ba8d47104a73707
-size 4899177160

 version https://git-lfs.github.com/spec/v1
+oid sha256:5843387d5c61880345dcb96dab915373df5a32908f0a9638bbe934d26c6e2605
+size 4899161900

model-00009-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:67fee5c31901b59cf6017a11e3fb2c15b25c333b33dfec3fe90d65f61354e2a2
-size 4899143244

 version https://git-lfs.github.com/spec/v1
+oid sha256:b0cce284a3fc84b43351440fc4a09103469bf0bfbef10d6aa0c548acbbdaa5d7
+size 4899161900

model-00010-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:c9980ef9ec00e1aa090e923da5c8947b9b902d71777e2610b0340613b8a6a4ee
-size 4983044820

 version https://git-lfs.github.com/spec/v1
+oid sha256:6b9ac707f663b09db5853b037fba63e8e7163c3cbc914ef4c9857572dea636ee
+size 4983047080

model-00011-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:cd28304dd794e6df10364c1048776e7af227d170bc86803985017b2b31ae62ab
-size 4899126640

 version https://git-lfs.github.com/spec/v1
+oid sha256:e7969b5046cd09601629a89b2264208a1f95ca0bba766331b68869a7d99d12f2
+size 4899146040

model-00012-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:151dd6e67a528f542ebccca6a811c03bbc13b2079373729d08c2e0ac56035625
-size 4899159724

 version https://git-lfs.github.com/spec/v1
+oid sha256:43786ea2bcb24de175b1c56ccb692f07d8759e616d1e66072be7b2be3d1789c7
+size 4899145404

model-00013-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:e741ab629c5a87d89122fde5b724bd2a110453df9a58ad2701f505c5925f1c8d
-size 4899159724

 version https://git-lfs.github.com/spec/v1
+oid sha256:e06841c6be0dced0db1961b53c76619f53b5fee463525adc8427d67d1f85bc15
+size 4899161908

model-00014-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:63408a09613132b527a1f4292151080c5fbc9d46bd0a25ae4772a162cdfb3097
-size 4983045008

 version https://git-lfs.github.com/spec/v1
+oid sha256:e49c43d251ef8d3945971db6e0e06f4505c2e2fb18293c1803ac4f457fad588f
+size 4983046852

model-00015-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:ec8872fdb2c24a07fb930f4eb0434647225e4ae432af373c8fdfedef4be6b9c6
-size 4899143864

 version https://git-lfs.github.com/spec/v1
+oid sha256:7e6f13ab0b56631f154111e2a11b875a8f6fb37d1c41f595bd330bd568f1ce40
+size 4899128768

model-00016-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:9b593cd8b74dc66ed7eed45303164f8f9285d15caded44ce42b0b746cacf890a
-size 4899143244

 version https://git-lfs.github.com/spec/v1
+oid sha256:da9a88b14df4918329ef19a89938a3a444477281af915101b483508391af5c5f
+size 4899161900

model-00017-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:0a98024b3b3c5b5815babc4dbd03972fa105f60dac44e9cde9cc7ef4376e1a93
-size 4899159732

 version https://git-lfs.github.com/spec/v1
+oid sha256:1e68b8b71dfb0dcf9222275e8db0f69d02ee21b03bc6e7e4fec6427ad8ca82b9
+size 4899179400

model-00018-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:52aa8f6213e3c18cd489dd723bfcf9a574a319d77317171644b4113fcaf253de
-size 4127387820

 version https://git-lfs.github.com/spec/v1
+oid sha256:7cde51ac7c61430a36ae18caac5f0b5491bb005f6f2af21807ec3ff1540454f8
+size 4983030364

model-00019-of-00019.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:c845d7ccba30e2e0eb8ab20aa313d5ed45dca6634e214a62ed323b70462fa828
-size 2101346432

 version https://git-lfs.github.com/spec/v1
+oid sha256:9f13ee3f4daeec8770d7996f5af786be664c899fdf6fc854ea89a00214bc9683
+size 3082882248

model.safetensors.index.json CHANGED Viewed

The diff for this file is too large to render. See raw diff

preprocessor_config.json ADDED Viewed

	@@ -0,0 +1,25 @@

+{
+  "do_convert_rgb": true,
+  "do_normalize": true,
+  "do_pad": true,
+  "do_rescale": true,
+  "do_resize": true,
+  "image_mean": [
+    0.48145466,
+    0.4578275,
+    0.40821073
+  ],
+  "image_processor_type": "MllamaImageProcessor",
+  "image_std": [
+    0.26862954,
+    0.26130258,
+    0.27577711
+  ],
+  "max_image_tiles": 4,
+  "resample": 2,
+  "rescale_factor": 0.00392156862745098,
+  "size": {
+    "height": 560,
+    "width": 560
+  }
+}

special_tokens_map.json CHANGED Viewed

@@ -13,11 +13,5 @@
     "rstrip": false,
     "single_word": false
   },
-  "pad_token": {
-    "content": "<|finetune_right_pad_id|>",
-    "lstrip": false,
-    "normalized": false,
-    "rstrip": false,
-    "single_word": false
-  }
 }

     "rstrip": false,
     "single_word": false
   },
+  "pad_token": "<|eot_id|>"
 }

tokenizer.json CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:47be6519609d58a5f29b3497045b8a2798d0d0978955ea90a893ad80e2ecdd4d
-size 17208880

 version https://git-lfs.github.com/spec/v1
+oid sha256:2950f10d620c6db4032082fe810ee88b29d8cda2caabfbcbf30f40eec988741c
+size 17210350

tokenizer_config.json CHANGED Viewed

@@ -2065,7 +2065,8 @@
     "input_ids",
     "attention_mask"
   ],
-  "model_max_length": 131072,
-  "pad_token": "<|finetune_right_pad_id|>",
   "tokenizer_class": "PreTrainedTokenizerFast"
 }

     "input_ids",
     "attention_mask"
   ],
+  "model_max_length": 512,
+  "pad_token": "<|eot_id|>",
+  "padding_side": "left",
   "tokenizer_class": "PreTrainedTokenizerFast"
 }