ashnrk commited on Jul 11, 2023

Commit

763ed4d

•

1 Parent(s): d9e8317

End of training

Browse files

Files changed (49) hide show

README.md +17 -0
checkpoint-1000/optimizer.bin +3 -0
checkpoint-1000/pytorch_model.bin +3 -0
checkpoint-1000/random_states_0.pkl +3 -0
checkpoint-1000/scheduler.bin +3 -0
checkpoint-1500/optimizer.bin +3 -0
checkpoint-1500/pytorch_model.bin +3 -0
checkpoint-1500/random_states_0.pkl +3 -0
checkpoint-1500/scheduler.bin +3 -0
checkpoint-2000/optimizer.bin +3 -0
checkpoint-2000/pytorch_model.bin +3 -0
checkpoint-2000/random_states_0.pkl +3 -0
checkpoint-2000/scheduler.bin +3 -0
checkpoint-2500/optimizer.bin +3 -0
checkpoint-2500/pytorch_model.bin +3 -0
checkpoint-2500/random_states_0.pkl +3 -0
checkpoint-2500/scheduler.bin +3 -0
checkpoint-3000/optimizer.bin +3 -0
checkpoint-3000/pytorch_model.bin +3 -0
checkpoint-3000/random_states_0.pkl +3 -0
checkpoint-3000/scheduler.bin +3 -0
checkpoint-500/optimizer.bin +3 -0
checkpoint-500/pytorch_model.bin +3 -0
checkpoint-500/random_states_0.pkl +3 -0
checkpoint-500/scheduler.bin +3 -0
feature_extractor/preprocessor_config.json +28 -0
learned_embeds-steps-1000.bin +3 -0
learned_embeds-steps-1500.bin +3 -0
learned_embeds-steps-2000.bin +3 -0
learned_embeds-steps-2500.bin +3 -0
learned_embeds-steps-3000.bin +3 -0
learned_embeds-steps-500.bin +3 -0
learned_embeds.bin +3 -0
logs/textual_inversion/1689072905.693046/events.out.tfevents.1689072905.ip-172-31-25-234.1798668.1 +3 -0
logs/textual_inversion/1689072905.6944845/hparams.yml +46 -0
logs/textual_inversion/events.out.tfevents.1689072905.ip-172-31-25-234.1798668.0 +3 -0
model_index.json +33 -0
scheduler/scheduler_config.json +20 -0
text_encoder/config.json +25 -0
text_encoder/pytorch_model.bin +3 -0
tokenizer/added_tokens.json +3 -0
tokenizer/merges.txt +0 -0
tokenizer/special_tokens_map.json +24 -0
tokenizer/tokenizer_config.json +33 -0
tokenizer/vocab.json +0 -0
unet/config.json +67 -0
unet/diffusion_pytorch_model.bin +3 -0
vae/config.json +31 -0
vae/diffusion_pytorch_model.bin +3 -0

README.md ADDED Viewed

	@@ -0,0 +1,17 @@

+---
+license: creativeml-openrail-m
+base_model: stabilityai/stable-diffusion-2-1
+tags:
+- stable-diffusion
+- stable-diffusion-diffusers
+- text-to-image
+- diffusers
+- textual_inversion
+inference: true
+---
+# Textual inversion text2image fine-tuning - ashnrk/textual_inversion_perm_crop
+These are textual inversion adaption weights for stabilityai/stable-diffusion-2-1. You can find some example images in the following.

checkpoint-1000/optimizer.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:2c0d1450c67851c98eb06f0c1e81521d0e7617dbae8519ba3e40cb2484ec54da
+size 404760109

checkpoint-1000/pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:ab51ec247aa328ded5342793ebb2ab834f3fdb26a9edf0e6f7a2ce0ab2878a70
+size 1361701921

checkpoint-1000/random_states_0.pkl ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:8b3896b58dc0ed3a7722bacc9db183409265fb7b83a781b00bcd5c18c9864b14
+size 21731

checkpoint-1000/scheduler.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:e4fefee7f0cd3cd72549c0d880d30f20865cdb0b9e2ff7f6414504c4b2e9f0da
+size 563

checkpoint-1500/optimizer.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:56d5e355c8aacc2a5c82d9f8a995fb704475f0d579c07ff55b445f9b80355b08
+size 404760109

checkpoint-1500/pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:09d3db51e38d73ae793a68e82a2cef46fa4a00493644ec953bc527047d7dbe82
+size 1361701921

checkpoint-1500/random_states_0.pkl ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:a90c5aab54e020b11dcef62db9e05d79d5e8d3c64ea7a8654a4eb8347c93c7b8
+size 21795

checkpoint-1500/scheduler.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:844da8963017941e63785b9c97b8add92acc2cd25941c28fb340a858881df1bf
+size 563

checkpoint-2000/optimizer.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:e0a8ac78073611ecef0b6a95e76d238f52ffc6a507e77b8cf92ea337a8ccbc93
+size 404760109

checkpoint-2000/pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:df45ba374e467f49414ceb4858ad9404bdd53e9dd5038073c10a0ec37acf21a4
+size 1361701921

checkpoint-2000/random_states_0.pkl ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:994f02470986883505352297f54d6ce7a14dd4e945689fbe56612ef4fe9ae7d5
+size 21731

checkpoint-2000/scheduler.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:cbf4fc1f2d7f293314583514528a85b4410bc5064cfda7e58674f7e1d643273c
+size 563

checkpoint-2500/optimizer.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:3d7f03af363e034d44334faaf2a2bf348922418365ad01c018c72f638d8273fb
+size 404760109

checkpoint-2500/pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6c8b64287bf2168860a95de05b255670c07ab174b7500750e6f324176a7f759b
+size 1361701921

checkpoint-2500/random_states_0.pkl ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6f7086d80246559be65a60008630fd835633dafb11c2608386dca765851caf70
+size 21795

checkpoint-2500/scheduler.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:4740df90e3b20bbc1c987eb976fb0a841bb921549ef0b5da564eb30b6591fa90
+size 563

checkpoint-3000/optimizer.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:1ced3c404ac6eef9de332afa88f71143ed2b7696bb130950794f91381b0bd437
+size 404760109

checkpoint-3000/pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:c72d244c3c5c06d8d91c3b739ea2f1a79b82927656693abbc8279d5cd5a0de1f
+size 1361701921

checkpoint-3000/random_states_0.pkl ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:50fbb95969f35f7e1592c930c42e61d42ced51ba36c7336849dc748422a4d16d
+size 21795

checkpoint-3000/scheduler.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:40a514eb416e461c32b87bc59fb69f01b75bd71e657d33697170ae8eab25546d
+size 563

checkpoint-500/optimizer.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:547600509168de8236305fa49603bd966a5fa2852d4378b9d4b66b8ea086596f
+size 404760109

checkpoint-500/pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:806c7b61cdd37c594740b335709d3d14db24eeeccc0635fe4cc0110eb954125a
+size 1361701921

checkpoint-500/random_states_0.pkl ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6c583219302242893ed8033e5ffc8e60a9edecbbeee150cb6d999d0d96d2a8cc
+size 21795

checkpoint-500/scheduler.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:4033b05306d90c4d61604702df45c09b47067814076296e9dc54ee6e1e25ece5
+size 563

feature_extractor/preprocessor_config.json ADDED Viewed

	@@ -0,0 +1,28 @@

+{
+  "crop_size": {
+    "height": 224,
+    "width": 224
+  },
+  "do_center_crop": true,
+  "do_convert_rgb": true,
+  "do_normalize": true,
+  "do_rescale": true,
+  "do_resize": true,
+  "feature_extractor_type": "CLIPFeatureExtractor",
+  "image_mean": [
+    0.48145466,
+    0.4578275,
+    0.40821073
+  ],
+  "image_processor_type": "CLIPImageProcessor",
+  "image_std": [
+    0.26862954,
+    0.26130258,
+    0.27577711
+  ],
+  "resample": 3,
+  "rescale_factor": 0.00392156862745098,
+  "size": {
+    "shortest_edge": 224
+  }
+}

learned_embeds-steps-1000.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:382637a7d15512e29075632e63030528aa825cb4de7df1df92f0db3c37b9c22d
+size 5025

learned_embeds-steps-1500.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:e2d34f8d0130240f88f2c28e9b3b2836961db21b80ad86d0e52ac01cd4aa3039
+size 5025

learned_embeds-steps-2000.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:855f428b46c2b63a82b8755b42005953522b7559340938641638dcd2a2e64e43
+size 5025

learned_embeds-steps-2500.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:8cad3dee7a1be21e474c4aa757b226944439113fcc09e305cebc5c04f1a9a970
+size 5025

learned_embeds-steps-3000.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:cf5f90d04443ebfa0ef6c02cc9309f3b077e6de7900369ca1785378f3fead0c1
+size 5025

learned_embeds-steps-500.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:f57f5716f03afc9126d2a96bd31425e4409bba3a2907ca4b65e18da468d2ae5e
+size 5022

learned_embeds.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:63334efc8fe67431aa108500c7a716554c7d754b893f295c74b59e419909a390
+size 4864

logs/textual_inversion/1689072905.693046/events.out.tfevents.1689072905.ip-172-31-25-234.1798668.1 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:8f202fce608a51b5b47b5a19590d250711c176625c61554f14c332cbf6656b7a
+size 2325

logs/textual_inversion/1689072905.6944845/hparams.yml ADDED Viewed

	@@ -0,0 +1,46 @@

+adam_beta1: 0.9
+adam_beta2: 0.999
+adam_epsilon: 1.0e-08
+adam_weight_decay: 0.01
+allow_tf32: false
+center_crop: false
+checkpointing_steps: 500
+checkpoints_total_limit: null
+dataloader_num_workers: 0
+enable_xformers_memory_efficient_attention: true
+gradient_accumulation_steps: 4
+gradient_checkpointing: false
+hub_model_id: null
+hub_token: null
+initializer_token: farm
+learnable_property: object
+learning_rate: 0.016
+local_rank: 0
+logging_dir: logs
+lr_num_cycles: 1
+lr_scheduler: constant
+lr_warmup_steps: 0
+max_train_steps: 3000
+mixed_precision: 'no'
+num_train_epochs: 94
+num_validation_images: 4
+num_vectors: 1
+output_dir: textual_inversion_perm_crop
+placeholder_token: <perm-crop-sat>
+pretrained_model_name_or_path: stabilityai/stable-diffusion-2-1
+push_to_hub: true
+repeats: 100
+report_to: tensorboard
+resolution: 512
+resume_from_checkpoint: null
+revision: null
+save_as_full_pipeline: false
+save_steps: 500
+scale_lr: true
+seed: null
+tokenizer_name: null
+train_batch_size: 1
+train_data_dir: /home/ubuntu/datasets/Eurosat_few_shot_10/PermanentCrop
+validation_epochs: null
+validation_prompt: null
+validation_steps: 100

logs/textual_inversion/events.out.tfevents.1689072905.ip-172-31-25-234.1798668.0 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:92b58dfa81cb85c3f7231dc605e21de08ef8cea4ce93d8c8d59c5b7829857ab2
+size 983642

model_index.json ADDED Viewed

	@@ -0,0 +1,33 @@

+{
+  "_class_name": "StableDiffusionPipeline",
+  "_diffusers_version": "0.18.0.dev0",
+  "feature_extractor": [
+    "transformers",
+    "CLIPImageProcessor"
+  ],
+  "requires_safety_checker": false,
+  "safety_checker": [
+    null,
+    null
+  ],
+  "scheduler": [
+    "diffusers",
+    "DDIMScheduler"
+  ],
+  "text_encoder": [
+    "transformers",
+    "CLIPTextModel"
+  ],
+  "tokenizer": [
+    "transformers",
+    "CLIPTokenizer"
+  ],
+  "unet": [
+    "diffusers",
+    "UNet2DConditionModel"
+  ],
+  "vae": [
+    "diffusers",
+    "AutoencoderKL"
+  ]
+}

scheduler/scheduler_config.json ADDED Viewed

	@@ -0,0 +1,20 @@

+{
+  "_class_name": "DDIMScheduler",
+  "_diffusers_version": "0.18.0.dev0",
+  "beta_end": 0.012,
+  "beta_schedule": "scaled_linear",
+  "beta_start": 0.00085,
+  "clip_sample": false,
+  "clip_sample_range": 1.0,
+  "dynamic_thresholding_ratio": 0.995,
+  "num_train_timesteps": 1000,
+  "prediction_type": "v_prediction",
+  "rescale_betas_zero_snr": false,
+  "sample_max_value": 1.0,
+  "set_alpha_to_one": false,
+  "skip_prk_steps": true,
+  "steps_offset": 1,
+  "thresholding": false,
+  "timestep_spacing": "leading",
+  "trained_betas": null
+}

text_encoder/config.json ADDED Viewed

	@@ -0,0 +1,25 @@

+{
+  "_name_or_path": "stabilityai/stable-diffusion-2-1",
+  "architectures": [
+    "CLIPTextModel"
+  ],
+  "attention_dropout": 0.0,
+  "bos_token_id": 0,
+  "dropout": 0.0,
+  "eos_token_id": 2,
+  "hidden_act": "gelu",
+  "hidden_size": 1024,
+  "initializer_factor": 1.0,
+  "initializer_range": 0.02,
+  "intermediate_size": 4096,
+  "layer_norm_eps": 1e-05,
+  "max_position_embeddings": 77,
+  "model_type": "clip_text_model",
+  "num_attention_heads": 16,
+  "num_hidden_layers": 23,
+  "pad_token_id": 1,
+  "projection_dim": 512,
+  "torch_dtype": "float32",
+  "transformers_version": "4.30.1",
+  "vocab_size": 49409
+}

text_encoder/pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:eecfdd7a311f5684db0c1fa3ca8316aedf6544aa805cb5b3975e7061eafcd766
+size 1361684001

tokenizer/added_tokens.json ADDED Viewed

	@@ -0,0 +1,3 @@

+{
+  "<perm-crop-sat>": 49408
+}

tokenizer/merges.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

tokenizer/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,24 @@

+{
+  "bos_token": {
+    "content": "<|startoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  },
+  "eos_token": {
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": "!",
+  "unk_token": {
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  }
+}

tokenizer/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,33 @@

+{
+  "add_prefix_space": false,
+  "bos_token": {
+    "__type": "AddedToken",
+    "content": "<|startoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  },
+  "clean_up_tokenization_spaces": true,
+  "do_lower_case": true,
+  "eos_token": {
+    "__type": "AddedToken",
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  },
+  "errors": "replace",
+  "model_max_length": 77,
+  "pad_token": "<|endoftext|>",
+  "tokenizer_class": "CLIPTokenizer",
+  "unk_token": {
+    "__type": "AddedToken",
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  }
+}

tokenizer/vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff

unet/config.json ADDED Viewed

	@@ -0,0 +1,67 @@

+{
+  "_class_name": "UNet2DConditionModel",
+  "_diffusers_version": "0.18.0.dev0",
+  "_name_or_path": "stabilityai/stable-diffusion-2-1",
+  "act_fn": "silu",
+  "addition_embed_type": null,
+  "addition_embed_type_num_heads": 64,
+  "attention_head_dim": [
+    5,
+    10,
+    20,
+    20
+  ],
+  "block_out_channels": [
+    320,
+    640,
+    1280,
+    1280
+  ],
+  "center_input_sample": false,
+  "class_embed_type": null,
+  "class_embeddings_concat": false,
+  "conv_in_kernel": 3,
+  "conv_out_kernel": 3,
+  "cross_attention_dim": 1024,
+  "cross_attention_norm": null,
+  "down_block_types": [
+    "CrossAttnDownBlock2D",
+    "CrossAttnDownBlock2D",
+    "CrossAttnDownBlock2D",
+    "DownBlock2D"
+  ],
+  "downsample_padding": 1,
+  "dual_cross_attention": false,
+  "encoder_hid_dim": null,
+  "encoder_hid_dim_type": null,
+  "flip_sin_to_cos": true,
+  "freq_shift": 0,
+  "in_channels": 4,
+  "layers_per_block": 2,
+  "mid_block_only_cross_attention": null,
+  "mid_block_scale_factor": 1,
+  "mid_block_type": "UNetMidBlock2DCrossAttn",
+  "norm_eps": 1e-05,
+  "norm_num_groups": 32,
+  "num_class_embeds": null,
+  "only_cross_attention": false,
+  "out_channels": 4,
+  "projection_class_embeddings_input_dim": null,
+  "resnet_out_scale_factor": 1.0,
+  "resnet_skip_time_act": false,
+  "resnet_time_scale_shift": "default",
+  "sample_size": 96,
+  "time_cond_proj_dim": null,
+  "time_embedding_act_fn": null,
+  "time_embedding_dim": null,
+  "time_embedding_type": "positional",
+  "timestep_post_act": null,
+  "up_block_types": [
+    "UpBlock2D",
+    "CrossAttnUpBlock2D",
+    "CrossAttnUpBlock2D",
+    "CrossAttnUpBlock2D"
+  ],
+  "upcast_attention": true,
+  "use_linear_projection": true
+}

unet/diffusion_pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:29c7f15709062e790deddc2c6d3d6a3e0a99d4a533b6bfd187bdea64a5e7a961
+size 3463934693

vae/config.json ADDED Viewed

	@@ -0,0 +1,31 @@

+{
+  "_class_name": "AutoencoderKL",
+  "_diffusers_version": "0.18.0.dev0",
+  "_name_or_path": "stabilityai/stable-diffusion-2-1",
+  "act_fn": "silu",
+  "block_out_channels": [
+    128,
+    256,
+    512,
+    512
+  ],
+  "down_block_types": [
+    "DownEncoderBlock2D",
+    "DownEncoderBlock2D",
+    "DownEncoderBlock2D",
+    "DownEncoderBlock2D"
+  ],
+  "in_channels": 3,
+  "latent_channels": 4,
+  "layers_per_block": 2,
+  "norm_num_groups": 32,
+  "out_channels": 3,
+  "sample_size": 768,
+  "scaling_factor": 0.18215,
+  "up_block_types": [
+    "UpDecoderBlock2D",
+    "UpDecoderBlock2D",
+    "UpDecoderBlock2D",
+    "UpDecoderBlock2D"
+  ]
+}

vae/diffusion_pytorch_model.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:185b0c03485b4048bb6158de087df301b79ea187844c76ae91cd4cda207282a2
+size 334715569