Chufan Wu commited on
Commit
96f5d38
·
1 Parent(s): c1fe24d
Files changed (31) hide show
  1. checkpoint-500/optimizer.bin +0 -3
  2. checkpoint-500/pytorch_lora_weights.safetensors +0 -3
  3. checkpoint-500/scaler.pt +0 -3
  4. checkpoint-500/scheduler.bin +0 -3
  5. checkpoint-1000/scheduler.bin → logs/text2image-fine-tune-sdxl/1694021617.6494288/events.out.tfevents.1694021617.twinkle1.794128.1 +2 -2
  6. logs/{text2image-fine-tune/1694016628.386469 → text2image-fine-tune-sdxl/1694021617.651228}/hparams.yml +15 -14
  7. checkpoint-500/random_states_0.pkl → logs/text2image-fine-tune-sdxl/events.out.tfevents.1694021617.twinkle1.794128.0 +2 -2
  8. logs/text2image-fine-tune/1694016628.3844774/events.out.tfevents.1694016628.twinkle1.765799.1 +0 -3
  9. logs/text2image-fine-tune/1694019947.4911246/events.out.tfevents.1694019947.twinkle1.781724.1 +0 -3
  10. logs/text2image-fine-tune/1694019947.4924138/hparams.yml +0 -51
  11. logs/text2image-fine-tune/events.out.tfevents.1694016628.twinkle1.765799.0 +0 -3
  12. logs/text2image-fine-tune/events.out.tfevents.1694019947.twinkle1.781724.0 +0 -3
  13. model_index.json +34 -0
  14. pytorch_lora_weights.safetensors +0 -3
  15. scheduler/scheduler_config.json +18 -0
  16. text_encoder/config.json +25 -0
  17. checkpoint-1000/optimizer.bin → text_encoder/model.safetensors +2 -2
  18. text_encoder_2/config.json +25 -0
  19. checkpoint-1000/random_states_0.pkl → text_encoder_2/model.safetensors +2 -2
  20. tokenizer/merges.txt +0 -0
  21. tokenizer/special_tokens_map.json +24 -0
  22. tokenizer/tokenizer_config.json +33 -0
  23. tokenizer/vocab.json +0 -0
  24. tokenizer_2/merges.txt +0 -0
  25. tokenizer_2/special_tokens_map.json +24 -0
  26. tokenizer_2/tokenizer_config.json +33 -0
  27. tokenizer_2/vocab.json +0 -0
  28. unet/config.json +72 -0
  29. checkpoint-1000/scaler.pt → unet/diffusion_pytorch_model.safetensors +2 -2
  30. vae/config.json +32 -0
  31. checkpoint-1000/pytorch_lora_weights.safetensors → vae/diffusion_pytorch_model.safetensors +2 -2
checkpoint-500/optimizer.bin DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:3e7497c67898dca80cc2acbd43ac726fb9fb9d67810a908a751b5ae882fc1b50
3
- size 47392445
 
 
 
 
checkpoint-500/pytorch_lora_weights.safetensors DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:359ae367f84b4dc4f59c5baa9cafee3774117cd1443b28b300846000349e6353
3
- size 23401064
 
 
 
 
checkpoint-500/scaler.pt DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:a3f196a54202bb4ba1220e8c59f42f9cda0702d68ea83147d814c2fb2f36b8f2
3
- size 557
 
 
 
 
checkpoint-500/scheduler.bin DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:62783ff6f0ae2c2ccc177b097e430c03940c5e70c38a191df273c5cf7ab1227c
3
- size 563
 
 
 
 
checkpoint-1000/scheduler.bin → logs/text2image-fine-tune-sdxl/1694021617.6494288/events.out.tfevents.1694021617.twinkle1.794128.1 RENAMED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:963c3085720a3a3ca1d178559a0f2dc5d5b36fa395a4663f0194bfac4a754038
3
- size 563
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:7b782dd43d192bc2b310904a36d9b30c0023fd43900f16f5ae532454aff7e71c
3
+ size 2403
logs/{text2image-fine-tune/1694016628.386469 → text2image-fine-tune-sdxl/1694021617.651228}/hparams.yml RENAMED
@@ -5,47 +5,48 @@ adam_weight_decay: 0.01
5
  allow_tf32: false
6
  cache_dir: null
7
  caption_column: text
8
- center_crop: false
9
- checkpointing_steps: 500
10
  checkpoints_total_limit: null
11
  dataloader_num_workers: 0
12
  dataset_config_name: null
13
  dataset_name: JAWCF/maps
14
- enable_xformers_memory_efficient_attention: false
15
- gradient_accumulation_steps: 1
16
- gradient_checkpointing: false
 
17
  hub_model_id: null
18
  hub_token: null
19
  image_column: image
20
- learning_rate: 0.0001
21
  local_rank: -1
22
  logging_dir: logs
23
  lr_scheduler: constant
24
  lr_warmup_steps: 0
25
  max_grad_norm: 1.0
26
  max_train_samples: null
27
- max_train_steps: 88
28
  mixed_precision: fp16
29
  noise_offset: 0
30
- num_train_epochs: 2
31
  num_validation_images: 4
32
- output_dir: sdxl-maps
33
  prediction_type: null
34
  pretrained_model_name_or_path: stabilityai/stable-diffusion-xl-base-1.0
35
  pretrained_vae_model_name_or_path: madebyollin/sdxl-vae-fp16-fix
 
36
  push_to_hub: false
37
  random_flip: true
38
- rank: 4
39
  report_to: tensorboard
40
- resolution: 1024
41
  resume_from_checkpoint: null
42
  revision: null
43
  scale_lr: false
44
- seed: 42
45
  snr_gamma: null
46
  train_batch_size: 1
47
  train_data_dir: null
48
- train_text_encoder: false
49
- use_8bit_adam: false
50
  validation_epochs: 1
51
  validation_prompt: null
 
5
  allow_tf32: false
6
  cache_dir: null
7
  caption_column: text
8
+ center_crop: true
9
+ checkpointing_steps: 1000
10
  checkpoints_total_limit: null
11
  dataloader_num_workers: 0
12
  dataset_config_name: null
13
  dataset_name: JAWCF/maps
14
+ enable_xformers_memory_efficient_attention: true
15
+ force_snr_gamma: false
16
+ gradient_accumulation_steps: 4
17
+ gradient_checkpointing: true
18
  hub_model_id: null
19
  hub_token: null
20
  image_column: image
21
+ learning_rate: 1.0e-06
22
  local_rank: -1
23
  logging_dir: logs
24
  lr_scheduler: constant
25
  lr_warmup_steps: 0
26
  max_grad_norm: 1.0
27
  max_train_samples: null
28
+ max_train_steps: 3000
29
  mixed_precision: fp16
30
  noise_offset: 0
31
+ num_train_epochs: 273
32
  num_validation_images: 4
33
+ output_dir: maps-xl
34
  prediction_type: null
35
  pretrained_model_name_or_path: stabilityai/stable-diffusion-xl-base-1.0
36
  pretrained_vae_model_name_or_path: madebyollin/sdxl-vae-fp16-fix
37
+ proportion_empty_prompts: 0.2
38
  push_to_hub: false
39
  random_flip: true
 
40
  report_to: tensorboard
41
+ resolution: 512
42
  resume_from_checkpoint: null
43
  revision: null
44
  scale_lr: false
45
+ seed: null
46
  snr_gamma: null
47
  train_batch_size: 1
48
  train_data_dir: null
49
+ use_8bit_adam: true
50
+ use_ema: false
51
  validation_epochs: 1
52
  validation_prompt: null
checkpoint-500/random_states_0.pkl → logs/text2image-fine-tune-sdxl/events.out.tfevents.1694021617.twinkle1.794128.0 RENAMED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:8f2bfff117f6074e4eb51eb6fd03c26fb5dead6a6052cd468bc56da4cc96a2e4
3
- size 14663
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:45cee1b83d8665760b1064fbc320335eef81f02c99aa559138604aba65323819
3
+ size 146961
logs/text2image-fine-tune/1694016628.3844774/events.out.tfevents.1694016628.twinkle1.765799.1 DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:ef6db7372ba09f2ddf869a4160d08611715cf27b629804052e5a55bde8cdef85
3
- size 2365
 
 
 
 
logs/text2image-fine-tune/1694019947.4911246/events.out.tfevents.1694019947.twinkle1.781724.1 DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:41af42758fab67a82ffd65267d30fae1e8664977451e7e5b76048b95eb138fd0
3
- size 2365
 
 
 
 
logs/text2image-fine-tune/1694019947.4924138/hparams.yml DELETED
@@ -1,51 +0,0 @@
1
- adam_beta1: 0.9
2
- adam_beta2: 0.999
3
- adam_epsilon: 1.0e-08
4
- adam_weight_decay: 0.01
5
- allow_tf32: false
6
- cache_dir: null
7
- caption_column: text
8
- center_crop: false
9
- checkpointing_steps: 500
10
- checkpoints_total_limit: null
11
- dataloader_num_workers: 0
12
- dataset_config_name: null
13
- dataset_name: JAWCF/maps
14
- enable_xformers_memory_efficient_attention: false
15
- gradient_accumulation_steps: 1
16
- gradient_checkpointing: false
17
- hub_model_id: null
18
- hub_token: null
19
- image_column: image
20
- learning_rate: 0.0001
21
- local_rank: -1
22
- logging_dir: logs
23
- lr_scheduler: constant
24
- lr_warmup_steps: 0
25
- max_grad_norm: 1.0
26
- max_train_samples: null
27
- max_train_steps: 1320
28
- mixed_precision: fp16
29
- noise_offset: 0
30
- num_train_epochs: 30
31
- num_validation_images: 4
32
- output_dir: sdxl-maps
33
- prediction_type: null
34
- pretrained_model_name_or_path: stabilityai/stable-diffusion-xl-base-1.0
35
- pretrained_vae_model_name_or_path: madebyollin/sdxl-vae-fp16-fix
36
- push_to_hub: false
37
- random_flip: true
38
- rank: 4
39
- report_to: tensorboard
40
- resolution: 1024
41
- resume_from_checkpoint: null
42
- revision: null
43
- scale_lr: false
44
- seed: 42
45
- snr_gamma: null
46
- train_batch_size: 1
47
- train_data_dir: null
48
- train_text_encoder: false
49
- use_8bit_adam: false
50
- validation_epochs: 1
51
- validation_prompt: null
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
logs/text2image-fine-tune/events.out.tfevents.1694016628.twinkle1.765799.0 DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:88167d9ca1f86c66d14d5c5e45d9932e15496d37132c503a6e3adcf34d8aee5b
3
- size 4312
 
 
 
 
logs/text2image-fine-tune/events.out.tfevents.1694019947.twinkle1.781724.0 DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:7b8f7475120f2ae9c1491b67ab47ec319f46e1c58741719ba815d00e745c07f8
3
- size 64641
 
 
 
 
model_index.json ADDED
@@ -0,0 +1,34 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_class_name": "StableDiffusionXLPipeline",
3
+ "_diffusers_version": "0.21.0.dev0",
4
+ "_name_or_path": "stabilityai/stable-diffusion-xl-base-1.0",
5
+ "force_zeros_for_empty_prompt": true,
6
+ "scheduler": [
7
+ "diffusers",
8
+ "EulerDiscreteScheduler"
9
+ ],
10
+ "text_encoder": [
11
+ "transformers",
12
+ "CLIPTextModel"
13
+ ],
14
+ "text_encoder_2": [
15
+ "transformers",
16
+ "CLIPTextModelWithProjection"
17
+ ],
18
+ "tokenizer": [
19
+ "transformers",
20
+ "CLIPTokenizer"
21
+ ],
22
+ "tokenizer_2": [
23
+ "transformers",
24
+ "CLIPTokenizer"
25
+ ],
26
+ "unet": [
27
+ "diffusers",
28
+ "UNet2DConditionModel"
29
+ ],
30
+ "vae": [
31
+ "diffusers",
32
+ "AutoencoderKL"
33
+ ]
34
+ }
pytorch_lora_weights.safetensors DELETED
@@ -1,3 +0,0 @@
1
- version https://git-lfs.github.com/spec/v1
2
- oid sha256:2c1e559a4f52e0513d0e73fde88397280d69c03588fa0671de694943ad553428
3
- size 23401064
 
 
 
 
scheduler/scheduler_config.json ADDED
@@ -0,0 +1,18 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_class_name": "EulerDiscreteScheduler",
3
+ "_diffusers_version": "0.21.0.dev0",
4
+ "beta_end": 0.012,
5
+ "beta_schedule": "scaled_linear",
6
+ "beta_start": 0.00085,
7
+ "clip_sample": false,
8
+ "interpolation_type": "linear",
9
+ "num_train_timesteps": 1000,
10
+ "prediction_type": "epsilon",
11
+ "sample_max_value": 1.0,
12
+ "set_alpha_to_one": false,
13
+ "skip_prk_steps": true,
14
+ "steps_offset": 1,
15
+ "timestep_spacing": "leading",
16
+ "trained_betas": null,
17
+ "use_karras_sigmas": false
18
+ }
text_encoder/config.json ADDED
@@ -0,0 +1,25 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "/home/wucf/.cache/huggingface/hub/models--stabilityai--stable-diffusion-xl-base-1.0/snapshots/f898a3e026e802f68796b95e9702464bac78d76f/text_encoder",
3
+ "architectures": [
4
+ "CLIPTextModel"
5
+ ],
6
+ "attention_dropout": 0.0,
7
+ "bos_token_id": 0,
8
+ "dropout": 0.0,
9
+ "eos_token_id": 2,
10
+ "hidden_act": "quick_gelu",
11
+ "hidden_size": 768,
12
+ "initializer_factor": 1.0,
13
+ "initializer_range": 0.02,
14
+ "intermediate_size": 3072,
15
+ "layer_norm_eps": 1e-05,
16
+ "max_position_embeddings": 77,
17
+ "model_type": "clip_text_model",
18
+ "num_attention_heads": 12,
19
+ "num_hidden_layers": 12,
20
+ "pad_token_id": 1,
21
+ "projection_dim": 768,
22
+ "torch_dtype": "float16",
23
+ "transformers_version": "4.33.0",
24
+ "vocab_size": 49408
25
+ }
checkpoint-1000/optimizer.bin → text_encoder/model.safetensors RENAMED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:2d535d837672c7ef94ddab64bdad01628a00d95428df06dd0020855f5a4ffb86
3
- size 47392445
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:660c6f5b1abae9dc498ac2d21e1347d2abdb0cf6c0c0c8576cd796491d9a6cdd
3
+ size 246144152
text_encoder_2/config.json ADDED
@@ -0,0 +1,25 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_name_or_path": "/home/wucf/.cache/huggingface/hub/models--stabilityai--stable-diffusion-xl-base-1.0/snapshots/f898a3e026e802f68796b95e9702464bac78d76f/text_encoder_2",
3
+ "architectures": [
4
+ "CLIPTextModelWithProjection"
5
+ ],
6
+ "attention_dropout": 0.0,
7
+ "bos_token_id": 0,
8
+ "dropout": 0.0,
9
+ "eos_token_id": 2,
10
+ "hidden_act": "gelu",
11
+ "hidden_size": 1280,
12
+ "initializer_factor": 1.0,
13
+ "initializer_range": 0.02,
14
+ "intermediate_size": 5120,
15
+ "layer_norm_eps": 1e-05,
16
+ "max_position_embeddings": 77,
17
+ "model_type": "clip_text_model",
18
+ "num_attention_heads": 20,
19
+ "num_hidden_layers": 32,
20
+ "pad_token_id": 1,
21
+ "projection_dim": 1280,
22
+ "torch_dtype": "float16",
23
+ "transformers_version": "4.33.0",
24
+ "vocab_size": 49408
25
+ }
checkpoint-1000/random_states_0.pkl → text_encoder_2/model.safetensors RENAMED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:192f79e73e3b59beb5b5695bc0a83081bf5cbb9f3e42eda7b23c62f73945d1e8
3
- size 14663
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:ec310df2af79c318e24d20511b601a591ca8cd4f1fce1d8dff822a356bcdb1f4
3
+ size 1389382176
tokenizer/merges.txt ADDED
The diff for this file is too large to render. See raw diff
 
tokenizer/special_tokens_map.json ADDED
@@ -0,0 +1,24 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "bos_token": {
3
+ "content": "<|startoftext|>",
4
+ "lstrip": false,
5
+ "normalized": true,
6
+ "rstrip": false,
7
+ "single_word": false
8
+ },
9
+ "eos_token": {
10
+ "content": "<|endoftext|>",
11
+ "lstrip": false,
12
+ "normalized": true,
13
+ "rstrip": false,
14
+ "single_word": false
15
+ },
16
+ "pad_token": "<|endoftext|>",
17
+ "unk_token": {
18
+ "content": "<|endoftext|>",
19
+ "lstrip": false,
20
+ "normalized": true,
21
+ "rstrip": false,
22
+ "single_word": false
23
+ }
24
+ }
tokenizer/tokenizer_config.json ADDED
@@ -0,0 +1,33 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "add_prefix_space": false,
3
+ "bos_token": {
4
+ "__type": "AddedToken",
5
+ "content": "<|startoftext|>",
6
+ "lstrip": false,
7
+ "normalized": true,
8
+ "rstrip": false,
9
+ "single_word": false
10
+ },
11
+ "clean_up_tokenization_spaces": true,
12
+ "do_lower_case": true,
13
+ "eos_token": {
14
+ "__type": "AddedToken",
15
+ "content": "<|endoftext|>",
16
+ "lstrip": false,
17
+ "normalized": true,
18
+ "rstrip": false,
19
+ "single_word": false
20
+ },
21
+ "errors": "replace",
22
+ "model_max_length": 77,
23
+ "pad_token": "<|endoftext|>",
24
+ "tokenizer_class": "CLIPTokenizer",
25
+ "unk_token": {
26
+ "__type": "AddedToken",
27
+ "content": "<|endoftext|>",
28
+ "lstrip": false,
29
+ "normalized": true,
30
+ "rstrip": false,
31
+ "single_word": false
32
+ }
33
+ }
tokenizer/vocab.json ADDED
The diff for this file is too large to render. See raw diff
 
tokenizer_2/merges.txt ADDED
The diff for this file is too large to render. See raw diff
 
tokenizer_2/special_tokens_map.json ADDED
@@ -0,0 +1,24 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "bos_token": {
3
+ "content": "<|startoftext|>",
4
+ "lstrip": false,
5
+ "normalized": true,
6
+ "rstrip": false,
7
+ "single_word": false
8
+ },
9
+ "eos_token": {
10
+ "content": "<|endoftext|>",
11
+ "lstrip": false,
12
+ "normalized": true,
13
+ "rstrip": false,
14
+ "single_word": false
15
+ },
16
+ "pad_token": "!",
17
+ "unk_token": {
18
+ "content": "<|endoftext|>",
19
+ "lstrip": false,
20
+ "normalized": true,
21
+ "rstrip": false,
22
+ "single_word": false
23
+ }
24
+ }
tokenizer_2/tokenizer_config.json ADDED
@@ -0,0 +1,33 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "add_prefix_space": false,
3
+ "bos_token": {
4
+ "__type": "AddedToken",
5
+ "content": "<|startoftext|>",
6
+ "lstrip": false,
7
+ "normalized": true,
8
+ "rstrip": false,
9
+ "single_word": false
10
+ },
11
+ "clean_up_tokenization_spaces": true,
12
+ "do_lower_case": true,
13
+ "eos_token": {
14
+ "__type": "AddedToken",
15
+ "content": "<|endoftext|>",
16
+ "lstrip": false,
17
+ "normalized": true,
18
+ "rstrip": false,
19
+ "single_word": false
20
+ },
21
+ "errors": "replace",
22
+ "model_max_length": 77,
23
+ "pad_token": "!",
24
+ "tokenizer_class": "CLIPTokenizer",
25
+ "unk_token": {
26
+ "__type": "AddedToken",
27
+ "content": "<|endoftext|>",
28
+ "lstrip": false,
29
+ "normalized": true,
30
+ "rstrip": false,
31
+ "single_word": false
32
+ }
33
+ }
tokenizer_2/vocab.json ADDED
The diff for this file is too large to render. See raw diff
 
unet/config.json ADDED
@@ -0,0 +1,72 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_class_name": "UNet2DConditionModel",
3
+ "_diffusers_version": "0.21.0.dev0",
4
+ "_name_or_path": "stabilityai/stable-diffusion-xl-base-1.0",
5
+ "act_fn": "silu",
6
+ "addition_embed_type": "text_time",
7
+ "addition_embed_type_num_heads": 64,
8
+ "addition_time_embed_dim": 256,
9
+ "attention_head_dim": [
10
+ 5,
11
+ 10,
12
+ 20
13
+ ],
14
+ "attention_type": "default",
15
+ "block_out_channels": [
16
+ 320,
17
+ 640,
18
+ 1280
19
+ ],
20
+ "center_input_sample": false,
21
+ "class_embed_type": null,
22
+ "class_embeddings_concat": false,
23
+ "conv_in_kernel": 3,
24
+ "conv_out_kernel": 3,
25
+ "cross_attention_dim": 2048,
26
+ "cross_attention_norm": null,
27
+ "down_block_types": [
28
+ "DownBlock2D",
29
+ "CrossAttnDownBlock2D",
30
+ "CrossAttnDownBlock2D"
31
+ ],
32
+ "downsample_padding": 1,
33
+ "dropout": 0.0,
34
+ "dual_cross_attention": false,
35
+ "encoder_hid_dim": null,
36
+ "encoder_hid_dim_type": null,
37
+ "flip_sin_to_cos": true,
38
+ "freq_shift": 0,
39
+ "in_channels": 4,
40
+ "layers_per_block": 2,
41
+ "mid_block_only_cross_attention": null,
42
+ "mid_block_scale_factor": 1,
43
+ "mid_block_type": "UNetMidBlock2DCrossAttn",
44
+ "norm_eps": 1e-05,
45
+ "norm_num_groups": 32,
46
+ "num_attention_heads": null,
47
+ "num_class_embeds": null,
48
+ "only_cross_attention": false,
49
+ "out_channels": 4,
50
+ "projection_class_embeddings_input_dim": 2816,
51
+ "resnet_out_scale_factor": 1.0,
52
+ "resnet_skip_time_act": false,
53
+ "resnet_time_scale_shift": "default",
54
+ "sample_size": 128,
55
+ "time_cond_proj_dim": null,
56
+ "time_embedding_act_fn": null,
57
+ "time_embedding_dim": null,
58
+ "time_embedding_type": "positional",
59
+ "timestep_post_act": null,
60
+ "transformer_layers_per_block": [
61
+ 1,
62
+ 2,
63
+ 10
64
+ ],
65
+ "up_block_types": [
66
+ "CrossAttnUpBlock2D",
67
+ "CrossAttnUpBlock2D",
68
+ "UpBlock2D"
69
+ ],
70
+ "upcast_attention": null,
71
+ "use_linear_projection": true
72
+ }
checkpoint-1000/scaler.pt → unet/diffusion_pytorch_model.safetensors RENAMED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:68cff80b680ddf6e7abbef98b5f336b97f9b5963e2209307f639383870e8cc71
3
- size 557
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c95444d15efd6410c6b5ff45db1d224f9d75d29476db9f0859bf47f5f2fa63ab
3
+ size 10270077736
vae/config.json ADDED
@@ -0,0 +1,32 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ {
2
+ "_class_name": "AutoencoderKL",
3
+ "_diffusers_version": "0.21.0.dev0",
4
+ "_name_or_path": "madebyollin/sdxl-vae-fp16-fix",
5
+ "act_fn": "silu",
6
+ "block_out_channels": [
7
+ 128,
8
+ 256,
9
+ 512,
10
+ 512
11
+ ],
12
+ "down_block_types": [
13
+ "DownEncoderBlock2D",
14
+ "DownEncoderBlock2D",
15
+ "DownEncoderBlock2D",
16
+ "DownEncoderBlock2D"
17
+ ],
18
+ "force_upcast": false,
19
+ "in_channels": 3,
20
+ "latent_channels": 4,
21
+ "layers_per_block": 2,
22
+ "norm_num_groups": 32,
23
+ "out_channels": 3,
24
+ "sample_size": 512,
25
+ "scaling_factor": 0.13025,
26
+ "up_block_types": [
27
+ "UpDecoderBlock2D",
28
+ "UpDecoderBlock2D",
29
+ "UpDecoderBlock2D",
30
+ "UpDecoderBlock2D"
31
+ ]
32
+ }
checkpoint-1000/pytorch_lora_weights.safetensors → vae/diffusion_pytorch_model.safetensors RENAMED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:f1a2c84c3cae8724cc199319bc8ced37ab0f144775f5c6dfd754c3fff263ebf9
3
- size 23401064
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:6353737672c94b96174cb590f711eac6edf2fcce5b6e91aa9d73c5adc589ee48
3
+ size 167335342