ArtChicken
/

vohwx_RealVisV4

Model card Files Files and versions Community

ArtChicken commited on Jun 12

Commit

e24e883

•

1 Parent(s): 26afc73

Upload 14 files

Browse files

Files changed (14) hide show

vohwx_RealVisV42024-06-12_11-26-22-save-585-15-0.safetensors +3 -0
vohwx_RealVisV42024-06-12_11-26-22-save-585-15-0.yaml +100 -0
vohwx_RealVisV42024-06-12_11-40-56-save-1170-30-0.safetensors +3 -0
vohwx_RealVisV42024-06-12_11-40-56-save-1170-30-0.yaml +100 -0
vohwx_RealVisV42024-06-12_11-55-50-save-1755-45-0.safetensors +3 -0
vohwx_RealVisV42024-06-12_11-55-50-save-1755-45-0.yaml +100 -0
vohwx_RealVisV42024-06-12_12-10-31-save-2340-60-0.safetensors +3 -0
vohwx_RealVisV42024-06-12_12-10-31-save-2340-60-0.yaml +100 -0
vohwx_RealVisV42024-06-12_12-25-12-save-2925-75-0.safetensors +3 -0
vohwx_RealVisV42024-06-12_12-25-12-save-2925-75-0.yaml +100 -0
vohwx_RealVisV42024-06-12_12-39-52-save-3510-90-0.safetensors +3 -0
vohwx_RealVisV42024-06-12_12-39-52-save-3510-90-0.yaml +100 -0
vohwx_RealVisV42024-06-12_12-54-43-save-4095-105-0.safetensors +3 -0
vohwx_RealVisV42024-06-12_12-54-43-save-4095-105-0.yaml +100 -0

vohwx_RealVisV42024-06-12_11-26-22-save-585-15-0.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:33b5c1f43cf56328290a695c59b4b95c4da02573ceade55ea97496d979d7ac7f
+size 6938084280

vohwx_RealVisV42024-06-12_11-26-22-save-585-15-0.yaml ADDED Viewed

	@@ -0,0 +1,100 @@

+model:
+  params:
+    conditioner_config:
+      params:
+        emb_models:
+        - input_key: txt
+          is_trainable: false
+          params:
+            layer: hidden
+            layer_idx: 11
+          target: sgm.modules.encoders.modules.FrozenCLIPEmbedder
+        - input_key: txt
+          is_trainable: false
+          params:
+            always_return_pooled: true
+            arch: ViT-bigG-14
+            freeze: true
+            layer: penultimate
+            legacy: false
+            version: laion2b_s39b_b160k
+          target: sgm.modules.encoders.modules.FrozenOpenCLIPEmbedder2
+        - input_key: original_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: crop_coords_top_left
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: target_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+      target: sgm.modules.GeneralConditioner
+    denoiser_config:
+      params:
+        discretization_config:
+          target: sgm.modules.diffusionmodules.discretizer.LegacyDDPMDiscretization
+        num_idx: 1000
+        scaling_config:
+          target: sgm.modules.diffusionmodules.denoiser_scaling.EpsScaling
+        weighting_config:
+          target: sgm.modules.diffusionmodules.denoiser_weighting.EpsWeighting
+      target: sgm.modules.diffusionmodules.denoiser.DiscreteDenoiser
+    disable_first_stage_autocast: true
+    first_stage_config:
+      params:
+        ddconfig:
+          attn_resolutions: []
+          attn_type: vanilla-xformers
+          ch: 128
+          ch_mult:
+          - 1
+          - 2
+          - 4
+          - 4
+          double_z: true
+          dropout: 0.0
+          in_channels: 3
+          num_res_blocks: 2
+          out_ch: 3
+          resolution: 256
+          z_channels: 4
+        embed_dim: 4
+        lossconfig:
+          target: torch.nn.Identity
+        monitor: val/rec_loss
+      target: sgm.models.autoencoder.AutoencoderKLInferenceWrapper
+    network_config:
+      params:
+        adm_in_channels: 2816
+        attention_resolutions:
+        - 4
+        - 2
+        channel_mult:
+        - 1
+        - 2
+        - 4
+        context_dim: 2048
+        in_channels: 4
+        legacy: false
+        model_channels: 320
+        num_classes: sequential
+        num_head_channels: 64
+        num_res_blocks: 2
+        out_channels: 4
+        spatial_transformer_attn_type: softmax-xformers
+        transformer_depth:
+        - 1
+        - 2
+        - 10
+        use_checkpoint: true
+        use_linear_in_transformer: true
+        use_spatial_transformer: true
+      target: sgm.modules.diffusionmodules.openaimodel.UNetModel
+    scale_factor: 0.13025
+  target: sgm.models.diffusion.DiffusionEngine

vohwx_RealVisV42024-06-12_11-40-56-save-1170-30-0.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:068aa4f0da0b5251ea3e5b8199ebfea04c16bfa236c7fe663efaa05b1a7c260c
+size 6938084280

vohwx_RealVisV42024-06-12_11-40-56-save-1170-30-0.yaml ADDED Viewed

	@@ -0,0 +1,100 @@

+model:
+  params:
+    conditioner_config:
+      params:
+        emb_models:
+        - input_key: txt
+          is_trainable: false
+          params:
+            layer: hidden
+            layer_idx: 11
+          target: sgm.modules.encoders.modules.FrozenCLIPEmbedder
+        - input_key: txt
+          is_trainable: false
+          params:
+            always_return_pooled: true
+            arch: ViT-bigG-14
+            freeze: true
+            layer: penultimate
+            legacy: false
+            version: laion2b_s39b_b160k
+          target: sgm.modules.encoders.modules.FrozenOpenCLIPEmbedder2
+        - input_key: original_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: crop_coords_top_left
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: target_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+      target: sgm.modules.GeneralConditioner
+    denoiser_config:
+      params:
+        discretization_config:
+          target: sgm.modules.diffusionmodules.discretizer.LegacyDDPMDiscretization
+        num_idx: 1000
+        scaling_config:
+          target: sgm.modules.diffusionmodules.denoiser_scaling.EpsScaling
+        weighting_config:
+          target: sgm.modules.diffusionmodules.denoiser_weighting.EpsWeighting
+      target: sgm.modules.diffusionmodules.denoiser.DiscreteDenoiser
+    disable_first_stage_autocast: true
+    first_stage_config:
+      params:
+        ddconfig:
+          attn_resolutions: []
+          attn_type: vanilla-xformers
+          ch: 128
+          ch_mult:
+          - 1
+          - 2
+          - 4
+          - 4
+          double_z: true
+          dropout: 0.0
+          in_channels: 3
+          num_res_blocks: 2
+          out_ch: 3
+          resolution: 256
+          z_channels: 4
+        embed_dim: 4
+        lossconfig:
+          target: torch.nn.Identity
+        monitor: val/rec_loss
+      target: sgm.models.autoencoder.AutoencoderKLInferenceWrapper
+    network_config:
+      params:
+        adm_in_channels: 2816
+        attention_resolutions:
+        - 4
+        - 2
+        channel_mult:
+        - 1
+        - 2
+        - 4
+        context_dim: 2048
+        in_channels: 4
+        legacy: false
+        model_channels: 320
+        num_classes: sequential
+        num_head_channels: 64
+        num_res_blocks: 2
+        out_channels: 4
+        spatial_transformer_attn_type: softmax-xformers
+        transformer_depth:
+        - 1
+        - 2
+        - 10
+        use_checkpoint: true
+        use_linear_in_transformer: true
+        use_spatial_transformer: true
+      target: sgm.modules.diffusionmodules.openaimodel.UNetModel
+    scale_factor: 0.13025
+  target: sgm.models.diffusion.DiffusionEngine

vohwx_RealVisV42024-06-12_11-55-50-save-1755-45-0.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6a97eefbb20cb0084bd00a7bf405982be4e7fa8932aa925358ae551e16b88269
+size 6938084280

vohwx_RealVisV42024-06-12_11-55-50-save-1755-45-0.yaml ADDED Viewed

	@@ -0,0 +1,100 @@

+model:
+  params:
+    conditioner_config:
+      params:
+        emb_models:
+        - input_key: txt
+          is_trainable: false
+          params:
+            layer: hidden
+            layer_idx: 11
+          target: sgm.modules.encoders.modules.FrozenCLIPEmbedder
+        - input_key: txt
+          is_trainable: false
+          params:
+            always_return_pooled: true
+            arch: ViT-bigG-14
+            freeze: true
+            layer: penultimate
+            legacy: false
+            version: laion2b_s39b_b160k
+          target: sgm.modules.encoders.modules.FrozenOpenCLIPEmbedder2
+        - input_key: original_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: crop_coords_top_left
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: target_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+      target: sgm.modules.GeneralConditioner
+    denoiser_config:
+      params:
+        discretization_config:
+          target: sgm.modules.diffusionmodules.discretizer.LegacyDDPMDiscretization
+        num_idx: 1000
+        scaling_config:
+          target: sgm.modules.diffusionmodules.denoiser_scaling.EpsScaling
+        weighting_config:
+          target: sgm.modules.diffusionmodules.denoiser_weighting.EpsWeighting
+      target: sgm.modules.diffusionmodules.denoiser.DiscreteDenoiser
+    disable_first_stage_autocast: true
+    first_stage_config:
+      params:
+        ddconfig:
+          attn_resolutions: []
+          attn_type: vanilla-xformers
+          ch: 128
+          ch_mult:
+          - 1
+          - 2
+          - 4
+          - 4
+          double_z: true
+          dropout: 0.0
+          in_channels: 3
+          num_res_blocks: 2
+          out_ch: 3
+          resolution: 256
+          z_channels: 4
+        embed_dim: 4
+        lossconfig:
+          target: torch.nn.Identity
+        monitor: val/rec_loss
+      target: sgm.models.autoencoder.AutoencoderKLInferenceWrapper
+    network_config:
+      params:
+        adm_in_channels: 2816
+        attention_resolutions:
+        - 4
+        - 2
+        channel_mult:
+        - 1
+        - 2
+        - 4
+        context_dim: 2048
+        in_channels: 4
+        legacy: false
+        model_channels: 320
+        num_classes: sequential
+        num_head_channels: 64
+        num_res_blocks: 2
+        out_channels: 4
+        spatial_transformer_attn_type: softmax-xformers
+        transformer_depth:
+        - 1
+        - 2
+        - 10
+        use_checkpoint: true
+        use_linear_in_transformer: true
+        use_spatial_transformer: true
+      target: sgm.modules.diffusionmodules.openaimodel.UNetModel
+    scale_factor: 0.13025
+  target: sgm.models.diffusion.DiffusionEngine

vohwx_RealVisV42024-06-12_12-10-31-save-2340-60-0.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:dd220ac685f18e889177992697e8e7c2ab43c984963de40c7b1892cd66900cf7
+size 6938084280

vohwx_RealVisV42024-06-12_12-10-31-save-2340-60-0.yaml ADDED Viewed

	@@ -0,0 +1,100 @@

+model:
+  params:
+    conditioner_config:
+      params:
+        emb_models:
+        - input_key: txt
+          is_trainable: false
+          params:
+            layer: hidden
+            layer_idx: 11
+          target: sgm.modules.encoders.modules.FrozenCLIPEmbedder
+        - input_key: txt
+          is_trainable: false
+          params:
+            always_return_pooled: true
+            arch: ViT-bigG-14
+            freeze: true
+            layer: penultimate
+            legacy: false
+            version: laion2b_s39b_b160k
+          target: sgm.modules.encoders.modules.FrozenOpenCLIPEmbedder2
+        - input_key: original_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: crop_coords_top_left
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: target_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+      target: sgm.modules.GeneralConditioner
+    denoiser_config:
+      params:
+        discretization_config:
+          target: sgm.modules.diffusionmodules.discretizer.LegacyDDPMDiscretization
+        num_idx: 1000
+        scaling_config:
+          target: sgm.modules.diffusionmodules.denoiser_scaling.EpsScaling
+        weighting_config:
+          target: sgm.modules.diffusionmodules.denoiser_weighting.EpsWeighting
+      target: sgm.modules.diffusionmodules.denoiser.DiscreteDenoiser
+    disable_first_stage_autocast: true
+    first_stage_config:
+      params:
+        ddconfig:
+          attn_resolutions: []
+          attn_type: vanilla-xformers
+          ch: 128
+          ch_mult:
+          - 1
+          - 2
+          - 4
+          - 4
+          double_z: true
+          dropout: 0.0
+          in_channels: 3
+          num_res_blocks: 2
+          out_ch: 3
+          resolution: 256
+          z_channels: 4
+        embed_dim: 4
+        lossconfig:
+          target: torch.nn.Identity
+        monitor: val/rec_loss
+      target: sgm.models.autoencoder.AutoencoderKLInferenceWrapper
+    network_config:
+      params:
+        adm_in_channels: 2816
+        attention_resolutions:
+        - 4
+        - 2
+        channel_mult:
+        - 1
+        - 2
+        - 4
+        context_dim: 2048
+        in_channels: 4
+        legacy: false
+        model_channels: 320
+        num_classes: sequential
+        num_head_channels: 64
+        num_res_blocks: 2
+        out_channels: 4
+        spatial_transformer_attn_type: softmax-xformers
+        transformer_depth:
+        - 1
+        - 2
+        - 10
+        use_checkpoint: true
+        use_linear_in_transformer: true
+        use_spatial_transformer: true
+      target: sgm.modules.diffusionmodules.openaimodel.UNetModel
+    scale_factor: 0.13025
+  target: sgm.models.diffusion.DiffusionEngine

vohwx_RealVisV42024-06-12_12-25-12-save-2925-75-0.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:f179525bd660a604a4a9fcec978efb699112b7896f53acb340cbfb813298ae2f
+size 6938084280

vohwx_RealVisV42024-06-12_12-25-12-save-2925-75-0.yaml ADDED Viewed

	@@ -0,0 +1,100 @@

+model:
+  params:
+    conditioner_config:
+      params:
+        emb_models:
+        - input_key: txt
+          is_trainable: false
+          params:
+            layer: hidden
+            layer_idx: 11
+          target: sgm.modules.encoders.modules.FrozenCLIPEmbedder
+        - input_key: txt
+          is_trainable: false
+          params:
+            always_return_pooled: true
+            arch: ViT-bigG-14
+            freeze: true
+            layer: penultimate
+            legacy: false
+            version: laion2b_s39b_b160k
+          target: sgm.modules.encoders.modules.FrozenOpenCLIPEmbedder2
+        - input_key: original_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: crop_coords_top_left
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: target_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+      target: sgm.modules.GeneralConditioner
+    denoiser_config:
+      params:
+        discretization_config:
+          target: sgm.modules.diffusionmodules.discretizer.LegacyDDPMDiscretization
+        num_idx: 1000
+        scaling_config:
+          target: sgm.modules.diffusionmodules.denoiser_scaling.EpsScaling
+        weighting_config:
+          target: sgm.modules.diffusionmodules.denoiser_weighting.EpsWeighting
+      target: sgm.modules.diffusionmodules.denoiser.DiscreteDenoiser
+    disable_first_stage_autocast: true
+    first_stage_config:
+      params:
+        ddconfig:
+          attn_resolutions: []
+          attn_type: vanilla-xformers
+          ch: 128
+          ch_mult:
+          - 1
+          - 2
+          - 4
+          - 4
+          double_z: true
+          dropout: 0.0
+          in_channels: 3
+          num_res_blocks: 2
+          out_ch: 3
+          resolution: 256
+          z_channels: 4
+        embed_dim: 4
+        lossconfig:
+          target: torch.nn.Identity
+        monitor: val/rec_loss
+      target: sgm.models.autoencoder.AutoencoderKLInferenceWrapper
+    network_config:
+      params:
+        adm_in_channels: 2816
+        attention_resolutions:
+        - 4
+        - 2
+        channel_mult:
+        - 1
+        - 2
+        - 4
+        context_dim: 2048
+        in_channels: 4
+        legacy: false
+        model_channels: 320
+        num_classes: sequential
+        num_head_channels: 64
+        num_res_blocks: 2
+        out_channels: 4
+        spatial_transformer_attn_type: softmax-xformers
+        transformer_depth:
+        - 1
+        - 2
+        - 10
+        use_checkpoint: true
+        use_linear_in_transformer: true
+        use_spatial_transformer: true
+      target: sgm.modules.diffusionmodules.openaimodel.UNetModel
+    scale_factor: 0.13025
+  target: sgm.models.diffusion.DiffusionEngine

vohwx_RealVisV42024-06-12_12-39-52-save-3510-90-0.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:ee81c86177a184f2815a231310359f6e3f4c79e3ba7103ad1dd8cdb877b52e8f
+size 6938084280

vohwx_RealVisV42024-06-12_12-39-52-save-3510-90-0.yaml ADDED Viewed

	@@ -0,0 +1,100 @@

+model:
+  params:
+    conditioner_config:
+      params:
+        emb_models:
+        - input_key: txt
+          is_trainable: false
+          params:
+            layer: hidden
+            layer_idx: 11
+          target: sgm.modules.encoders.modules.FrozenCLIPEmbedder
+        - input_key: txt
+          is_trainable: false
+          params:
+            always_return_pooled: true
+            arch: ViT-bigG-14
+            freeze: true
+            layer: penultimate
+            legacy: false
+            version: laion2b_s39b_b160k
+          target: sgm.modules.encoders.modules.FrozenOpenCLIPEmbedder2
+        - input_key: original_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: crop_coords_top_left
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: target_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+      target: sgm.modules.GeneralConditioner
+    denoiser_config:
+      params:
+        discretization_config:
+          target: sgm.modules.diffusionmodules.discretizer.LegacyDDPMDiscretization
+        num_idx: 1000
+        scaling_config:
+          target: sgm.modules.diffusionmodules.denoiser_scaling.EpsScaling
+        weighting_config:
+          target: sgm.modules.diffusionmodules.denoiser_weighting.EpsWeighting
+      target: sgm.modules.diffusionmodules.denoiser.DiscreteDenoiser
+    disable_first_stage_autocast: true
+    first_stage_config:
+      params:
+        ddconfig:
+          attn_resolutions: []
+          attn_type: vanilla-xformers
+          ch: 128
+          ch_mult:
+          - 1
+          - 2
+          - 4
+          - 4
+          double_z: true
+          dropout: 0.0
+          in_channels: 3
+          num_res_blocks: 2
+          out_ch: 3
+          resolution: 256
+          z_channels: 4
+        embed_dim: 4
+        lossconfig:
+          target: torch.nn.Identity
+        monitor: val/rec_loss
+      target: sgm.models.autoencoder.AutoencoderKLInferenceWrapper
+    network_config:
+      params:
+        adm_in_channels: 2816
+        attention_resolutions:
+        - 4
+        - 2
+        channel_mult:
+        - 1
+        - 2
+        - 4
+        context_dim: 2048
+        in_channels: 4
+        legacy: false
+        model_channels: 320
+        num_classes: sequential
+        num_head_channels: 64
+        num_res_blocks: 2
+        out_channels: 4
+        spatial_transformer_attn_type: softmax-xformers
+        transformer_depth:
+        - 1
+        - 2
+        - 10
+        use_checkpoint: true
+        use_linear_in_transformer: true
+        use_spatial_transformer: true
+      target: sgm.modules.diffusionmodules.openaimodel.UNetModel
+    scale_factor: 0.13025
+  target: sgm.models.diffusion.DiffusionEngine

vohwx_RealVisV42024-06-12_12-54-43-save-4095-105-0.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:3e7baa32a670ae75d742cd5d51d74145f3995d07734a2c53a68df216765622b2
+size 6938084280

vohwx_RealVisV42024-06-12_12-54-43-save-4095-105-0.yaml ADDED Viewed

	@@ -0,0 +1,100 @@

+model:
+  params:
+    conditioner_config:
+      params:
+        emb_models:
+        - input_key: txt
+          is_trainable: false
+          params:
+            layer: hidden
+            layer_idx: 11
+          target: sgm.modules.encoders.modules.FrozenCLIPEmbedder
+        - input_key: txt
+          is_trainable: false
+          params:
+            always_return_pooled: true
+            arch: ViT-bigG-14
+            freeze: true
+            layer: penultimate
+            legacy: false
+            version: laion2b_s39b_b160k
+          target: sgm.modules.encoders.modules.FrozenOpenCLIPEmbedder2
+        - input_key: original_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: crop_coords_top_left
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+        - input_key: target_size_as_tuple
+          is_trainable: false
+          params:
+            outdim: 256
+          target: sgm.modules.encoders.modules.ConcatTimestepEmbedderND
+      target: sgm.modules.GeneralConditioner
+    denoiser_config:
+      params:
+        discretization_config:
+          target: sgm.modules.diffusionmodules.discretizer.LegacyDDPMDiscretization
+        num_idx: 1000
+        scaling_config:
+          target: sgm.modules.diffusionmodules.denoiser_scaling.EpsScaling
+        weighting_config:
+          target: sgm.modules.diffusionmodules.denoiser_weighting.EpsWeighting
+      target: sgm.modules.diffusionmodules.denoiser.DiscreteDenoiser
+    disable_first_stage_autocast: true
+    first_stage_config:
+      params:
+        ddconfig:
+          attn_resolutions: []
+          attn_type: vanilla-xformers
+          ch: 128
+          ch_mult:
+          - 1
+          - 2
+          - 4
+          - 4
+          double_z: true
+          dropout: 0.0
+          in_channels: 3
+          num_res_blocks: 2
+          out_ch: 3
+          resolution: 256
+          z_channels: 4
+        embed_dim: 4
+        lossconfig:
+          target: torch.nn.Identity
+        monitor: val/rec_loss
+      target: sgm.models.autoencoder.AutoencoderKLInferenceWrapper
+    network_config:
+      params:
+        adm_in_channels: 2816
+        attention_resolutions:
+        - 4
+        - 2
+        channel_mult:
+        - 1
+        - 2
+        - 4
+        context_dim: 2048
+        in_channels: 4
+        legacy: false
+        model_channels: 320
+        num_classes: sequential
+        num_head_channels: 64
+        num_res_blocks: 2
+        out_channels: 4
+        spatial_transformer_attn_type: softmax-xformers
+        transformer_depth:
+        - 1
+        - 2
+        - 10
+        use_checkpoint: true
+        use_linear_in_transformer: true
+        use_spatial_transformer: true
+      target: sgm.modules.diffusionmodules.openaimodel.UNetModel
+    scale_factor: 0.13025
+  target: sgm.models.diffusion.DiffusionEngine