Upload 16 files

Browse files

Files changed (16) hide show

scheduler/scheduler_config.json +30 -0
text_encoder/config.json +63 -0
text_encoder/generation_config.json +14 -0
text_encoder/model.safetensors +3 -0
tokenizer/chat_template.json +4 -0
tokenizer/merges.txt +0 -0
tokenizer/preprocessor_config.json +21 -0
tokenizer/tokenizer.json +0 -0
tokenizer/tokenizer_config.json +239 -0
tokenizer/video_preprocessor_config.json +21 -0
tokenizer/vocab.json +0 -0
transformer/config.json +27 -0
transformer/diffusion_pytorch_model-00001-of-00001.safetensors +3 -0
transformer/diffusion_pytorch_model.safetensors.index.json +546 -0
vae/config.json +89 -0
vae/diffusion_pytorch_model.safetensors +3 -0

scheduler/scheduler_config.json ADDED Viewed

	@@ -0,0 +1,30 @@

+{
+  "_class_name": "DPMSolverMultistepScheduler",
+  "_diffusers_version": "0.33.1",
+  "algorithm_type": "dpmsolver++",
+  "beta_end": 0.02,
+  "beta_schedule": "linear",
+  "beta_start": 0.0001,
+  "dynamic_thresholding_ratio": 0.995,
+  "euler_at_final": false,
+  "final_sigmas_type": "zero",
+  "flow_shift": 3.0,
+  "lambda_min_clipped": -Infinity,
+  "lower_order_final": true,
+  "num_train_timesteps": 1000,
+  "prediction_type": "flow_prediction",
+  "rescale_betas_zero_snr": false,
+  "sample_max_value": 1.0,
+  "solver_order": 2,
+  "solver_type": "midpoint",
+  "steps_offset": 0,
+  "thresholding": false,
+  "timestep_spacing": "linspace",
+  "trained_betas": null,
+  "use_beta_sigmas": false,
+  "use_exponential_sigmas": false,
+  "use_flow_sigmas": true,
+  "use_karras_sigmas": false,
+  "use_lu_lambdas": false,
+  "variance_type": null
+}

text_encoder/config.json ADDED Viewed

	@@ -0,0 +1,63 @@

+{
+  "architectures": [
+    "Qwen3VLForConditionalGeneration"
+  ],
+  "image_token_id": 151655,
+  "model_type": "qwen3_vl",
+  "text_config": {
+    "attention_bias": false,
+    "attention_dropout": 0.0,
+    "bos_token_id": 151643,
+    "dtype": "bfloat16",
+    "eos_token_id": 151645,
+    "head_dim": 128,
+    "hidden_act": "silu",
+    "hidden_size": 2048,
+    "initializer_range": 0.02,
+    "intermediate_size": 6144,
+    "max_position_embeddings": 262144,
+    "model_type": "qwen3_vl_text",
+    "num_attention_heads": 16,
+    "num_hidden_layers": 28,
+    "num_key_value_heads": 8,
+    "rms_norm_eps": 1e-06,
+    "rope_scaling": {
+      "mrope_interleaved": true,
+      "mrope_section": [
+        24,
+        20,
+        20
+      ],
+      "rope_type": "default"
+    },
+    "rope_theta": 5000000,
+    "tie_word_embeddings": true,
+    "use_cache": true,
+    "vocab_size": 151936
+  },
+  "tie_word_embeddings": true,
+  "transformers_version": "4.57.1",
+  "video_token_id": 151656,
+  "vision_config": {
+    "deepstack_visual_indexes": [
+      5,
+      11,
+      17
+    ],
+    "depth": 24,
+    "hidden_act": "gelu_pytorch_tanh",
+    "hidden_size": 1024,
+    "in_channels": 3,
+    "initializer_range": 0.02,
+    "intermediate_size": 4096,
+    "model_type": "qwen3_vl",
+    "num_heads": 16,
+    "num_position_embeddings": 2304,
+    "out_hidden_size": 2048,
+    "patch_size": 16,
+    "spatial_merge_size": 2,
+    "temporal_patch_size": 2
+  },
+  "vision_end_token_id": 151653,
+  "vision_start_token_id": 151652
+}

text_encoder/generation_config.json ADDED Viewed

	@@ -0,0 +1,14 @@

+{
+    "bos_token_id": 151643,
+    "pad_token_id": 151643,
+    "do_sample": true,
+    "eos_token_id": [
+        151645,
+        151643
+    ],
+    "top_p": 0.8,
+    "top_k": 20,
+    "temperature": 0.7,
+    "repetition_penalty": 1.0,
+    "transformers_version": "4.56.0"
+}

text_encoder/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:7de1838c87a5349b016c26a1c3f7d2bc400a3d485f95ef39a7059ffd734977a0
+size 4255140312

tokenizer/chat_template.json ADDED Viewed

	@@ -0,0 +1,4 @@

+{
+    "chat_template": "{%- if tools %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if messages[0].role == 'system' %}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n\\n' }}\n    {%- endif %}\n    {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n    {%- for tool in tools %}\n        {{- \"\\n\" }}\n        {{- tool | tojson }}\n    {%- endfor %}\n    {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call><|im_end|>\\n\" }}\n{%- else %}\n    {%- if messages[0].role == 'system' %}\n        {{- '<|im_start|>system\\n' }}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- endif %}\n{%- endif %}\n{%- set image_count = namespace(value=0) %}\n{%- set video_count = namespace(value=0) %}\n{%- for message in messages %}\n    {%- if message.role == \"user\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content_item in message.content %}\n                {%- if 'text' in content_item %}\n                    {{- content_item.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and message.content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n{%- endif %}\n"
+}

tokenizer/merges.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

tokenizer/preprocessor_config.json ADDED Viewed

	@@ -0,0 +1,21 @@

+{
+    "size": {
+        "longest_edge": 16777216,
+        "shortest_edge": 65536
+    },
+    "patch_size": 16,
+    "temporal_patch_size": 2,
+    "merge_size": 2,
+    "image_mean": [
+        0.5,
+        0.5,
+        0.5
+    ],
+    "image_std": [
+        0.5,
+        0.5,
+        0.5
+    ],
+    "processor_class": "Qwen3VLProcessor",
+    "image_processor_type": "Qwen2VLImageProcessorFast"
+}

tokenizer/tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

tokenizer/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,239 @@

+{
+  "add_bos_token": false,
+  "add_prefix_space": false,
+  "added_tokens_decoder": {
+    "151643": {
+      "content": "<|endoftext|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151644": {
+      "content": "<|im_start|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151645": {
+      "content": "<|im_end|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151646": {
+      "content": "<|object_ref_start|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151647": {
+      "content": "<|object_ref_end|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151648": {
+      "content": "<|box_start|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151649": {
+      "content": "<|box_end|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151650": {
+      "content": "<|quad_start|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151651": {
+      "content": "<|quad_end|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151652": {
+      "content": "<|vision_start|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151653": {
+      "content": "<|vision_end|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151654": {
+      "content": "<|vision_pad|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151655": {
+      "content": "<|image_pad|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151656": {
+      "content": "<|video_pad|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151657": {
+      "content": "<tool_call>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151658": {
+      "content": "</tool_call>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151659": {
+      "content": "<|fim_prefix|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151660": {
+      "content": "<|fim_middle|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151661": {
+      "content": "<|fim_suffix|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151662": {
+      "content": "<|fim_pad|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151663": {
+      "content": "<|repo_name|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151664": {
+      "content": "<|file_sep|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151665": {
+      "content": "<tool_response>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151666": {
+      "content": "</tool_response>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151667": {
+      "content": "<think>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    },
+    "151668": {
+      "content": "</think>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": false
+    }
+  },
+  "additional_special_tokens": [
+    "<|im_start|>",
+    "<|im_end|>",
+    "<|object_ref_start|>",
+    "<|object_ref_end|>",
+    "<|box_start|>",
+    "<|box_end|>",
+    "<|quad_start|>",
+    "<|quad_end|>",
+    "<|vision_start|>",
+    "<|vision_end|>",
+    "<|vision_pad|>",
+    "<|image_pad|>",
+    "<|video_pad|>"
+  ],
+  "bos_token": null,
+  "chat_template": "{%- if tools %}\n    {{- '<|im_start|>system\\n' }}\n    {%- if messages[0].role == 'system' %}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n\\n' }}\n    {%- endif %}\n    {{- \"# Tools\\n\\nYou may call one or more functions to assist with the user query.\\n\\nYou are provided with function signatures within <tools></tools> XML tags:\\n<tools>\" }}\n    {%- for tool in tools %}\n        {{- \"\\n\" }}\n        {{- tool | tojson }}\n    {%- endfor %}\n    {{- \"\\n</tools>\\n\\nFor each function call, return a json object with function name and arguments within <tool_call></tool_call> XML tags:\\n<tool_call>\\n{\\\"name\\\": <function-name>, \\\"arguments\\\": <args-json-object>}\\n</tool_call><|im_end|>\\n\" }}\n{%- else %}\n    {%- if messages[0].role == 'system' %}\n        {{- '<|im_start|>system\\n' }}\n        {%- if messages[0].content is string %}\n            {{- messages[0].content }}\n        {%- else %}\n            {%- for content in messages[0].content %}\n                {%- if 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- endif %}\n{%- endif %}\n{%- set image_count = namespace(value=0) %}\n{%- set video_count = namespace(value=0) %}\n{%- for message in messages %}\n    {%- if message.role == \"user\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"assistant\" %}\n        {{- '<|im_start|>' + message.role + '\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content_item in message.content %}\n                {%- if 'text' in content_item %}\n                    {{- content_item.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {%- if message.tool_calls %}\n            {%- for tool_call in message.tool_calls %}\n                {%- if (loop.first and message.content) or (not loop.first) %}\n                    {{- '\\n' }}\n                {%- endif %}\n                {%- if tool_call.function %}\n                    {%- set tool_call = tool_call.function %}\n                {%- endif %}\n                {{- '<tool_call>\\n{\"name\": \"' }}\n                {{- tool_call.name }}\n                {{- '\", \"arguments\": ' }}\n                {%- if tool_call.arguments is string %}\n                    {{- tool_call.arguments }}\n                {%- else %}\n                    {{- tool_call.arguments | tojson }}\n                {%- endif %}\n                {{- '}\\n</tool_call>' }}\n            {%- endfor %}\n        {%- endif %}\n        {{- '<|im_end|>\\n' }}\n    {%- elif message.role == \"tool\" %}\n        {%- if loop.first or (messages[loop.index0 - 1].role != \"tool\") %}\n            {{- '<|im_start|>user' }}\n        {%- endif %}\n        {{- '\\n<tool_response>\\n' }}\n        {%- if message.content is string %}\n            {{- message.content }}\n        {%- else %}\n            {%- for content in message.content %}\n                {%- if content.type == 'image' or 'image' in content or 'image_url' in content %}\n                    {%- set image_count.value = image_count.value + 1 %}\n                    {%- if add_vision_id %}Picture {{ image_count.value }}: {% endif -%}\n                    <|vision_start|><|image_pad|><|vision_end|>\n                {%- elif content.type == 'video' or 'video' in content %}\n                    {%- set video_count.value = video_count.value + 1 %}\n                    {%- if add_vision_id %}Video {{ video_count.value }}: {% endif -%}\n                    <|vision_start|><|video_pad|><|vision_end|>\n                {%- elif 'text' in content %}\n                    {{- content.text }}\n                {%- endif %}\n            {%- endfor %}\n        {%- endif %}\n        {{- '\\n</tool_response>' }}\n        {%- if loop.last or (messages[loop.index0 + 1].role != \"tool\") %}\n            {{- '<|im_end|>\\n' }}\n        {%- endif %}\n    {%- endif %}\n{%- endfor %}\n{%- if add_generation_prompt %}\n    {{- '<|im_start|>assistant\\n' }}\n{%- endif %}\n",
+  "clean_up_tokenization_spaces": false,
+  "eos_token": "<|im_end|>",
+  "errors": "replace",
+  "model_max_length": 262144,
+  "pad_token": "<|endoftext|>",
+  "split_special_tokens": false,
+  "tokenizer_class": "Qwen2Tokenizer",
+  "unk_token": null
+}

tokenizer/video_preprocessor_config.json ADDED Viewed

	@@ -0,0 +1,21 @@

+{
+    "size": {
+        "longest_edge": 25165824,
+        "shortest_edge": 4096
+    },
+    "patch_size": 16,
+    "temporal_patch_size": 2,
+    "merge_size": 2,
+    "image_mean": [
+        0.5,
+        0.5,
+        0.5
+    ],
+    "image_std": [
+        0.5,
+        0.5,
+        0.5
+    ],
+    "processor_class": "Qwen3VLProcessor",
+    "video_processor_type": "Qwen3VLVideoProcessor"
+}

tokenizer/vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff

transformer/config.json ADDED Viewed

	@@ -0,0 +1,27 @@

+{
+  "_class_name": "VIBESanaEditingModel",
+  "_diffusers_version": "0.33.1",
+  "attention_bias": false,
+  "attention_head_dim": 32,
+  "caption_channels": 2304,
+  "cross_attention_dim": 2240,
+  "cross_attention_head_dim": 112,
+  "dropout": 0.0,
+  "guidance_embeds": false,
+  "in_channels": 64,
+  "interpolation_scale": null,
+  "mlp_ratio": 2.5,
+  "norm_elementwise_affine": false,
+  "norm_eps": 1e-06,
+  "num_attention_heads": 70,
+  "num_cross_attention_heads": 20,
+  "num_layers": 20,
+  "out_channels": 32,
+  "patch_size": 1,
+  "qk_norm": "rms_norm_across_heads",
+  "edit_head_input_dim": 2048,
+  "edit_head_stacks_num": 4,
+  "num_meta_queries": 224,
+  "meta_queries_dim": 2048,
+  "input_condition_type": "channel_cat"
+}

transformer/diffusion_pytorch_model-00001-of-00001.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:ac81406ec9f673a120695f83db41745f1b7b0e8a0b1f45a8de66aa420fc52c54
+size 3765399232

transformer/diffusion_pytorch_model.safetensors.index.json ADDED Viewed

	@@ -0,0 +1,546 @@

+{
+  "metadata": {
+    "total_size": 3765336896
+  },
+  "weight_map": {
+    "caption_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "caption_projection.linear_1.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "caption_projection.linear_1.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "caption_projection.linear_2.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "caption_projection.linear_2.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.attn.in_proj_bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.attn.in_proj_weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.attn.out_proj.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.attn.out_proj.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.input_norm.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.input_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.mlp.fc1.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.mlp.fc1.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.mlp.fc2.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.mlp.fc2.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.mlp.fc3.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.mlp.fc3.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.output_norm.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.0.output_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.attn.in_proj_bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.attn.in_proj_weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.attn.out_proj.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.attn.out_proj.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.input_norm.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.input_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.mlp.fc1.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.mlp.fc1.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.mlp.fc2.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.mlp.fc2.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.mlp.fc3.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.mlp.fc3.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.output_norm.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.1.output_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.attn.in_proj_bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.attn.in_proj_weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.attn.out_proj.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.attn.out_proj.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.input_norm.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.input_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.mlp.fc1.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.mlp.fc1.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.mlp.fc2.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.mlp.fc2.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.mlp.fc3.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.mlp.fc3.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.output_norm.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.2.output_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.attn.in_proj_bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.attn.in_proj_weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.attn.out_proj.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.attn.out_proj.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.input_norm.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.input_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.mlp.fc1.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.mlp.fc1.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.mlp.fc2.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.mlp.fc2.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.mlp.fc3.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.mlp.fc3.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.output_norm.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.attn_blocks.3.output_norm.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.input_proj.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.input_proj.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.output_proj.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.output_proj.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.output_proj.1.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "edit_head.output_proj.1.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "meta_queries": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "patch_embed.proj.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "patch_embed.proj.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "proj_out.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "proj_out.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "time_embed.emb.timestep_embedder.linear_1.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "time_embed.emb.timestep_embedder.linear_1.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "time_embed.emb.timestep_embedder.linear_2.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "time_embed.emb.timestep_embedder.linear_2.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "time_embed.linear.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "time_embed.linear.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.0.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.1.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.10.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.11.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.12.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.13.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.14.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.15.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.16.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.17.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.18.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.19.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.2.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.3.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.4.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.5.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.6.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.7.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.8.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn1.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn1.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn1.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn1.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn1.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn1.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn1.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.norm_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.norm_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.to_k.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.to_k.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.to_out.0.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.to_out.0.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.to_q.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.to_q.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.to_v.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.attn2.to_v.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.ff.conv_depth.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.ff.conv_depth.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.ff.conv_inverted.bias": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.ff.conv_inverted.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.ff.conv_point.weight": "diffusion_pytorch_model-00001-of-00001.safetensors",
+    "transformer_blocks.9.scale_shift_table": "diffusion_pytorch_model-00001-of-00001.safetensors"
+  }
+}

vae/config.json ADDED Viewed

	@@ -0,0 +1,89 @@

+{
+  "_class_name": "AutoencoderDC",
+  "_diffusers_version": "0.33.1",
+  "_name_or_path": "mit-han-lab/dc-ae-f32c32-sana-1.1-diffusers",
+  "attention_head_dim": 32,
+  "decoder_act_fns": "silu",
+  "decoder_block_out_channels": [
+    128,
+    256,
+    512,
+    512,
+    1024,
+    1024
+  ],
+  "decoder_block_types": [
+    "ResBlock",
+    "ResBlock",
+    "ResBlock",
+    "EfficientViTBlock",
+    "EfficientViTBlock",
+    "EfficientViTBlock"
+  ],
+  "decoder_layers_per_block": [
+    3,
+    3,
+    3,
+    3,
+    3,
+    3
+  ],
+  "decoder_norm_types": "rms_norm",
+  "decoder_qkv_multiscales": [
+    [],
+    [],
+    [],
+    [
+      5
+    ],
+    [
+      5
+    ],
+    [
+      5
+    ]
+  ],
+  "downsample_block_type": "Conv",
+  "encoder_block_out_channels": [
+    128,
+    256,
+    512,
+    512,
+    1024,
+    1024
+  ],
+  "encoder_block_types": [
+    "ResBlock",
+    "ResBlock",
+    "ResBlock",
+    "EfficientViTBlock",
+    "EfficientViTBlock",
+    "EfficientViTBlock"
+  ],
+  "encoder_layers_per_block": [
+    2,
+    2,
+    2,
+    3,
+    3,
+    3
+  ],
+  "encoder_qkv_multiscales": [
+    [],
+    [],
+    [],
+    [
+      5
+    ],
+    [
+      5
+    ],
+    [
+      5
+    ]
+  ],
+  "in_channels": 3,
+  "latent_channels": 32,
+  "scaling_factor": 0.41407,
+  "upsample_block_type": "interpolate"
+}

vae/diffusion_pytorch_model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:dfd991d1b54ffabf22745c5885589d8f2a7bc59930d95d92bd741c4fc64454bb
+size 1249044836