update

2025-08-18 14:02:05 +05:30 · 2025-08-18 13:29:07 +05:30 · 2025-08-18 13:20:59 +05:30 · 2025-08-18 11:40:02 +05:30
69 changed files with 176 additions and 383 deletions
@@ -14,10 +14,6 @@

 # QwenImage

-<div class="flex flex-wrap space-x-1">
-  <img alt="LoRA" src="https://img.shields.io/badge/LoRA-d8b4fe?style=flat"/>
-</div>
-
 Qwen-Image from the Qwen team is an image generation foundation model in the Qwen series that achieves significant advances in complex text rendering and precise image editing. Experiments show strong general capabilities in both image generation and editing, with exceptional performance in text rendering, especially for Chinese.

 Qwen-Image comes in the following variants:
@@ -90,12 +86,6 @@ image.save("qwen_fewsteps.png")

 </details>

-<Tip>
-
-The `guidance_scale` parameter in the pipeline is there to support future guidance-distilled models when they come up. Note that passing `guidance_scale` to the pipeline is ineffective. To enable classifier-free guidance, please pass `true_cfg_scale` and `negative_prompt` (even an empty negative prompt like " ") should enable classifier-free guidance computations.
-
-</Tip>
-
 ## QwenImagePipeline

 [[autodoc]] QwenImagePipeline
@@ -333,8 +333,6 @@ The general rule of thumb to keep in mind when preparing inputs for the VACE pip

 - Wan 2.1 and 2.2 support using [LightX2V LoRAs](https://huggingface.co/Kijai/WanVideo_comfy/tree/main/Lightx2v) to speed up inference. Using them on Wan 2.2 is slightly more involed. Refer to [this code snippet](https://github.com/huggingface/diffusers/pull/12040#issuecomment-3144185272) to learn more.

- Wan 2.2 has two denoisers. By default, LoRAs are only loaded into the first denoiser. One can set `load_into_transformer_2=True` to load LoRAs into the second denoiser. Refer to [this](https://github.com/huggingface/diffusers/pull/12074#issue-3292620048) and [this](https://github.com/huggingface/diffusers/pull/12074#issuecomment-3155896144) examples to learn more.
-
 ## WanPipeline

 [[autodoc]] WanPipeline
@@ -90,7 +90,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -88,7 +88,7 @@ from diffusers.utils.import_utils import is_xformers_available


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -95,7 +95,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -61,7 +61,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -52,7 +52,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -60,7 +60,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -43,7 +43,8 @@ from diffusers.utils import BaseOutput, check_min_version


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")
+

 class MarigoldDepthOutput(BaseOutput):
    """
@@ -74,7 +74,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -67,7 +67,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -80,7 +80,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -73,7 +73,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -79,7 +79,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -61,7 +61,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -61,7 +61,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = logging.getLogger(__name__)

@@ -66,7 +66,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)
 if is_torch_npu_available():
@@ -62,7 +62,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -62,7 +62,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)
 if is_torch_npu_available():
@@ -64,7 +64,7 @@ from diffusers.utils.import_utils import is_xformers_available


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -19,9 +19,8 @@ cd diffusers
 pip install -e .
 ```

-Install the requirements in the `examples/dreambooth` folder as shown below.
+Then cd in the example folder and run
 ```bash
-cd examples/dreambooth
 pip install -r requirements.txt
 ```

@@ -75,9 +75,9 @@ Now, we can launch training using:
 ```bash
 export MODEL_NAME="Qwen/Qwen-Image"
 export INSTANCE_DIR="dog"
-export OUTPUT_DIR="trained-qwenimage-lora"
+export OUTPUT_DIR="trained-sana-lora"

-accelerate launch train_dreambooth_lora_qwenimage.py \
+accelerate launch train_dreambooth_lora_sana.py \
  --pretrained_model_name_or_path=$MODEL_NAME  \
  --instance_data_dir=$INSTANCE_DIR \
  --output_dir=$OUTPUT_DIR \
@@ -64,7 +64,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -35,7 +35,7 @@ from diffusers.utils import check_min_version


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 # Cache compiled models across invocations of this script.
 cc.initialize_cache(os.path.expanduser("~/.cache/jax/compilation_cache"))
@@ -80,7 +80,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -75,7 +75,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -87,7 +87,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -73,7 +73,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.34.0.dev0")

 logger = get_logger(__name__)

@@ -75,7 +75,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -73,7 +73,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -75,7 +75,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -87,7 +87,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -73,7 +73,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -80,7 +80,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -64,7 +64,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -55,7 +55,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -58,7 +58,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -58,7 +58,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -60,7 +60,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -53,7 +53,7 @@ if is_wandb_available():


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -46,7 +46,7 @@ from diffusers.utils import check_min_version, is_wandb_available


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -46,7 +46,7 @@ from diffusers.utils import check_min_version, is_wandb_available


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -52,7 +52,7 @@ if is_wandb_available():


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -61,7 +61,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -57,7 +57,7 @@ if is_wandb_available():


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -49,7 +49,7 @@ from diffusers.utils import check_min_version


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = logging.getLogger(__name__)

@@ -56,7 +56,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -68,7 +68,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)
 if is_torch_npu_available():
@@ -55,7 +55,7 @@ from diffusers.utils.torch_utils import is_compiled_module


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)
 if is_torch_npu_available():
@@ -82,7 +82,7 @@ else:


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -56,7 +56,7 @@ else:
 # ------------------------------------------------------------------------------

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = logging.getLogger(__name__)

@@ -77,7 +77,7 @@ else:


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__)

@@ -29,7 +29,7 @@ from diffusers.utils.import_utils import is_xformers_available


 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -50,7 +50,7 @@ if is_wandb_available():
    import wandb

 # Will error if the minimal version of diffusers is not installed. Remove at your own risks.
-check_min_version("0.35.0")
+check_min_version("0.35.0.dev0")

 logger = get_logger(__name__, log_level="INFO")

@@ -269,7 +269,7 @@ version_range_max = max(sys.version_info[1], 10) + 1

 setup(
    name="diffusers",
-    version="0.35.1",  # expected format is one of x.y.z.dev0, or x.y.z.rc1 or x.y.z (no to dashes, yes to dots)
+    version="0.35.0.dev0",  # expected format is one of x.y.z.dev0, or x.y.z.rc1 or x.y.z (no to dashes, yes to dots)
    description="State-of-the-art diffusion in PyTorch and JAX.",
    long_description=open("README.md", "r", encoding="utf-8").read(),
    long_description_content_type="text/markdown",
@@ -1,4 +1,4 @@
-__version__ = "0.35.1"
+__version__ = "0.35.0.dev0"

 from typing import TYPE_CHECKING

@@ -754,11 +754,7 @@ class LoraBaseMixin:
        # Decompose weights into weights for denoiser and text encoders.
        _component_adapter_weights = {}
        for component in self._lora_loadable_modules:
-            model = getattr(self, component, None)
-            # To guard for cases like Wan. In Wan2.1 and WanVace, we have a single denoiser.
-            # Whereas in Wan 2.2, we have two denoisers.
-            if model is None:
-                continue
+            model = getattr(self, component)

            for adapter_name, weights in zip(adapter_names, adapter_weights):
                if isinstance(weights, dict):
@@ -1833,17 +1833,6 @@ def _convert_non_diffusers_wan_lora_to_diffusers(state_dict):
        k.startswith("time_projection") and k.endswith(".weight") for k in original_state_dict
    )

-    def get_alpha_scales(down_weight, alpha_key):
-        rank = down_weight.shape[0]
-        alpha = original_state_dict.pop(alpha_key).item()
-        scale = alpha / rank  # LoRA is scaled by 'alpha / rank' in forward pass, so we need to scale it back here
-        scale_down = scale
-        scale_up = 1.0
-        while scale_down * 2 < scale_up:
-            scale_down *= 2
-            scale_up /= 2
-        return scale_down, scale_up
-
    for key in list(original_state_dict.keys()):
        if key.endswith((".diff", ".diff_b")) and "norm" in key:
            # NOTE: we don't support this because norm layer diff keys are just zeroed values. We can support it
@@ -1863,26 +1852,15 @@ def _convert_non_diffusers_wan_lora_to_diffusers(state_dict):
    for i in range(min_block, max_block + 1):
        # Self-attention
        for o, c in zip(["q", "k", "v", "o"], ["to_q", "to_k", "to_v", "to_out.0"]):
-            alpha_key = f"blocks.{i}.self_attn.{o}.alpha"
-            has_alpha = alpha_key in original_state_dict
-            original_key_A = f"blocks.{i}.self_attn.{o}.{lora_down_key}.weight"
-            converted_key_A = f"blocks.{i}.attn1.{c}.lora_A.weight"
+            original_key = f"blocks.{i}.self_attn.{o}.{lora_down_key}.weight"
+            converted_key = f"blocks.{i}.attn1.{c}.lora_A.weight"
+            if original_key in original_state_dict:
+                converted_state_dict[converted_key] = original_state_dict.pop(original_key)

-            original_key_B = f"blocks.{i}.self_attn.{o}.{lora_up_key}.weight"
-            converted_key_B = f"blocks.{i}.attn1.{c}.lora_B.weight"
-
-            if has_alpha:
-                down_weight = original_state_dict.pop(original_key_A)
-                up_weight = original_state_dict.pop(original_key_B)
-                scale_down, scale_up = get_alpha_scales(down_weight, alpha_key)
-                converted_state_dict[converted_key_A] = down_weight * scale_down
-                converted_state_dict[converted_key_B] = up_weight * scale_up
-
-            else:
-                if original_key_A in original_state_dict:
-                    converted_state_dict[converted_key_A] = original_state_dict.pop(original_key_A)
-                if original_key_B in original_state_dict:
-                    converted_state_dict[converted_key_B] = original_state_dict.pop(original_key_B)
+            original_key = f"blocks.{i}.self_attn.{o}.{lora_up_key}.weight"
+            converted_key = f"blocks.{i}.attn1.{c}.lora_B.weight"
+            if original_key in original_state_dict:
+                converted_state_dict[converted_key] = original_state_dict.pop(original_key)

            original_key = f"blocks.{i}.self_attn.{o}.diff_b"
            converted_key = f"blocks.{i}.attn1.{c}.lora_B.bias"
@@ -1891,24 +1869,15 @@ def _convert_non_diffusers_wan_lora_to_diffusers(state_dict):

        # Cross-attention
        for o, c in zip(["q", "k", "v", "o"], ["to_q", "to_k", "to_v", "to_out.0"]):
-            alpha_key = f"blocks.{i}.cross_attn.{o}.alpha"
-            has_alpha = alpha_key in original_state_dict
-            original_key_A = f"blocks.{i}.cross_attn.{o}.{lora_down_key}.weight"
-            converted_key_A = f"blocks.{i}.attn2.{c}.lora_A.weight"
+            original_key = f"blocks.{i}.cross_attn.{o}.{lora_down_key}.weight"
+            converted_key = f"blocks.{i}.attn2.{c}.lora_A.weight"
+            if original_key in original_state_dict:
+                converted_state_dict[converted_key] = original_state_dict.pop(original_key)

-            original_key_B = f"blocks.{i}.cross_attn.{o}.{lora_up_key}.weight"
-            converted_key_B = f"blocks.{i}.attn2.{c}.lora_B.weight"
-
-            if original_key_A in original_state_dict:
-                down_weight = original_state_dict.pop(original_key_A)
-                converted_state_dict[converted_key_A] = down_weight
-            if original_key_B in original_state_dict:
-                up_weight = original_state_dict.pop(original_key_B)
-                converted_state_dict[converted_key_B] = up_weight
-            if has_alpha:
-                scale_down, scale_up = get_alpha_scales(down_weight, alpha_key)
-                converted_state_dict[converted_key_A] *= scale_down
-                converted_state_dict[converted_key_B] *= scale_up
+            original_key = f"blocks.{i}.cross_attn.{o}.{lora_up_key}.weight"
+            converted_key = f"blocks.{i}.attn2.{c}.lora_B.weight"
+            if original_key in original_state_dict:
+                converted_state_dict[converted_key] = original_state_dict.pop(original_key)

            original_key = f"blocks.{i}.cross_attn.{o}.diff_b"
            converted_key = f"blocks.{i}.attn2.{c}.lora_B.bias"
@@ -1917,24 +1886,15 @@ def _convert_non_diffusers_wan_lora_to_diffusers(state_dict):

        if is_i2v_lora:
            for o, c in zip(["k_img", "v_img"], ["add_k_proj", "add_v_proj"]):
-                alpha_key = f"blocks.{i}.cross_attn.{o}.alpha"
-                has_alpha = alpha_key in original_state_dict
-                original_key_A = f"blocks.{i}.cross_attn.{o}.{lora_down_key}.weight"
-                converted_key_A = f"blocks.{i}.attn2.{c}.lora_A.weight"
+                original_key = f"blocks.{i}.cross_attn.{o}.{lora_down_key}.weight"
+                converted_key = f"blocks.{i}.attn2.{c}.lora_A.weight"
+                if original_key in original_state_dict:
+                    converted_state_dict[converted_key] = original_state_dict.pop(original_key)

-                original_key_B = f"blocks.{i}.cross_attn.{o}.{lora_up_key}.weight"
-                converted_key_B = f"blocks.{i}.attn2.{c}.lora_B.weight"
-
-                if original_key_A in original_state_dict:
-                    down_weight = original_state_dict.pop(original_key_A)
-                    converted_state_dict[converted_key_A] = down_weight
-                if original_key_B in original_state_dict:
-                    up_weight = original_state_dict.pop(original_key_B)
-                    converted_state_dict[converted_key_B] = up_weight
-                if has_alpha:
-                    scale_down, scale_up = get_alpha_scales(down_weight, alpha_key)
-                    converted_state_dict[converted_key_A] *= scale_down
-                    converted_state_dict[converted_key_B] *= scale_up
+                original_key = f"blocks.{i}.cross_attn.{o}.{lora_up_key}.weight"
+                converted_key = f"blocks.{i}.attn2.{c}.lora_B.weight"
+                if original_key in original_state_dict:
+                    converted_state_dict[converted_key] = original_state_dict.pop(original_key)

                original_key = f"blocks.{i}.cross_attn.{o}.diff_b"
                converted_key = f"blocks.{i}.attn2.{c}.lora_B.bias"
@@ -1943,24 +1903,15 @@ def _convert_non_diffusers_wan_lora_to_diffusers(state_dict):

        # FFN
        for o, c in zip(["ffn.0", "ffn.2"], ["net.0.proj", "net.2"]):
-            alpha_key = f"blocks.{i}.{o}.alpha"
-            has_alpha = alpha_key in original_state_dict
-            original_key_A = f"blocks.{i}.{o}.{lora_down_key}.weight"
-            converted_key_A = f"blocks.{i}.ffn.{c}.lora_A.weight"
+            original_key = f"blocks.{i}.{o}.{lora_down_key}.weight"
+            converted_key = f"blocks.{i}.ffn.{c}.lora_A.weight"
+            if original_key in original_state_dict:
+                converted_state_dict[converted_key] = original_state_dict.pop(original_key)

-            original_key_B = f"blocks.{i}.{o}.{lora_up_key}.weight"
-            converted_key_B = f"blocks.{i}.ffn.{c}.lora_B.weight"
-
-            if original_key_A in original_state_dict:
-                down_weight = original_state_dict.pop(original_key_A)
-                converted_state_dict[converted_key_A] = down_weight
-            if original_key_B in original_state_dict:
-                up_weight = original_state_dict.pop(original_key_B)
-                converted_state_dict[converted_key_B] = up_weight
-            if has_alpha:
-                scale_down, scale_up = get_alpha_scales(down_weight, alpha_key)
-                converted_state_dict[converted_key_A] *= scale_down
-                converted_state_dict[converted_key_B] *= scale_up
+            original_key = f"blocks.{i}.{o}.{lora_up_key}.weight"
+            converted_key = f"blocks.{i}.ffn.{c}.lora_B.weight"
+            if original_key in original_state_dict:
+                converted_state_dict[converted_key] = original_state_dict.pop(original_key)

            original_key = f"blocks.{i}.{o}.diff_b"
            converted_key = f"blocks.{i}.ffn.{c}.lora_B.bias"
@@ -2129,74 +2080,6 @@ def _convert_non_diffusers_ltxv_lora_to_diffusers(state_dict, non_diffusers_pref


 def _convert_non_diffusers_qwen_lora_to_diffusers(state_dict):
-    has_lora_unet = any(k.startswith("lora_unet_") for k in state_dict)
-    if has_lora_unet:
-        state_dict = {k.removeprefix("lora_unet_"): v for k, v in state_dict.items()}
-
-        def convert_key(key: str) -> str:
-            prefix = "transformer_blocks"
-            if "." in key:
-                base, suffix = key.rsplit(".", 1)
-            else:
-                base, suffix = key, ""
-
-            start = f"{prefix}_"
-            rest = base[len(start) :]
-
-            if "." in rest:
-                head, tail = rest.split(".", 1)
-                tail = "." + tail
-            else:
-                head, tail = rest, ""
-
-            # Protected n-grams that must keep their internal underscores
-            protected = {
-                # pairs
-                ("to", "q"),
-                ("to", "k"),
-                ("to", "v"),
-                ("to", "out"),
-                ("add", "q"),
-                ("add", "k"),
-                ("add", "v"),
-                ("txt", "mlp"),
-                ("img", "mlp"),
-                ("txt", "mod"),
-                ("img", "mod"),
-                # triplets
-                ("add", "q", "proj"),
-                ("add", "k", "proj"),
-                ("add", "v", "proj"),
-                ("to", "add", "out"),
-            }
-
-            prot_by_len = {}
-            for ng in protected:
-                prot_by_len.setdefault(len(ng), set()).add(ng)
-
-            parts = head.split("_")
-            merged = []
-            i = 0
-            lengths_desc = sorted(prot_by_len.keys(), reverse=True)
-
-            while i < len(parts):
-                matched = False
-                for L in lengths_desc:
-                    if i + L <= len(parts) and tuple(parts[i : i + L]) in prot_by_len[L]:
-                        merged.append("_".join(parts[i : i + L]))
-                        i += L
-                        matched = True
-                        break
-                if not matched:
-                    merged.append(parts[i])
-                    i += 1
-
-            head_converted = ".".join(merged)
-            converted_base = f"{prefix}.{head_converted}{tail}"
-            return converted_base + (("." + suffix) if suffix else "")
-
-        state_dict = {convert_key(k): v for k, v in state_dict.items()}
-
    converted_state_dict = {}
    all_keys = list(state_dict.keys())
    down_key = ".lora_down.weight"
@@ -5065,7 +5065,7 @@ class WanLoraLoaderMixin(LoraBaseMixin):
    Load LoRA layers into [`WanTransformer3DModel`]. Specific to [`WanPipeline`] and `[WanImageToVideoPipeline`].
    """

-    _lora_loadable_modules = ["transformer", "transformer_2"]
+    _lora_loadable_modules = ["transformer"]
    transformer_name = TRANSFORMER_NAME

    @classmethod
@@ -5270,35 +5270,15 @@ class WanLoraLoaderMixin(LoraBaseMixin):
        if not is_correct_format:
            raise ValueError("Invalid LoRA checkpoint.")

-        load_into_transformer_2 = kwargs.pop("load_into_transformer_2", False)
-        if load_into_transformer_2:
-            if not hasattr(self, "transformer_2"):
-                raise AttributeError(
-                    f"'{type(self).__name__}' object has no attribute transformer_2"
-                    "Note that Wan2.1 models do not have a transformer_2 component."
-                    "Ensure the model has a transformer_2 component before setting load_into_transformer_2=True."
-                )
-            self.load_lora_into_transformer(
-                state_dict,
-                transformer=self.transformer_2,
-                adapter_name=adapter_name,
-                metadata=metadata,
-                _pipeline=self,
-                low_cpu_mem_usage=low_cpu_mem_usage,
-                hotswap=hotswap,
-            )
-        else:
-            self.load_lora_into_transformer(
-                state_dict,
-                transformer=getattr(self, self.transformer_name)
-                if not hasattr(self, "transformer")
-                else self.transformer,
-                adapter_name=adapter_name,
-                metadata=metadata,
-                _pipeline=self,
-                low_cpu_mem_usage=low_cpu_mem_usage,
-                hotswap=hotswap,
-            )
+        self.load_lora_into_transformer(
+            state_dict,
+            transformer=getattr(self, self.transformer_name) if not hasattr(self, "transformer") else self.transformer,
+            adapter_name=adapter_name,
+            metadata=metadata,
+            _pipeline=self,
+            low_cpu_mem_usage=low_cpu_mem_usage,
+            hotswap=hotswap,
+        )

    @classmethod
    # Copied from diffusers.loaders.lora_pipeline.SD3LoraLoaderMixin.load_lora_into_transformer with SD3Transformer2DModel->WanTransformer3DModel
@@ -5688,35 +5668,15 @@ class SkyReelsV2LoraLoaderMixin(LoraBaseMixin):
        if not is_correct_format:
            raise ValueError("Invalid LoRA checkpoint.")

-        load_into_transformer_2 = kwargs.pop("load_into_transformer_2", False)
-        if load_into_transformer_2:
-            if not hasattr(self, "transformer_2"):
-                raise AttributeError(
-                    f"'{type(self).__name__}' object has no attribute transformer_2"
-                    "Note that Wan2.1 models do not have a transformer_2 component."
-                    "Ensure the model has a transformer_2 component before setting load_into_transformer_2=True."
-                )
-            self.load_lora_into_transformer(
-                state_dict,
-                transformer=self.transformer_2,
-                adapter_name=adapter_name,
-                metadata=metadata,
-                _pipeline=self,
-                low_cpu_mem_usage=low_cpu_mem_usage,
-                hotswap=hotswap,
-            )
-        else:
-            self.load_lora_into_transformer(
-                state_dict,
-                transformer=getattr(self, self.transformer_name)
-                if not hasattr(self, "transformer")
-                else self.transformer,
-                adapter_name=adapter_name,
-                metadata=metadata,
-                _pipeline=self,
-                low_cpu_mem_usage=low_cpu_mem_usage,
-                hotswap=hotswap,
-            )
+        self.load_lora_into_transformer(
+            state_dict,
+            transformer=getattr(self, self.transformer_name) if not hasattr(self, "transformer") else self.transformer,
+            adapter_name=adapter_name,
+            metadata=metadata,
+            _pipeline=self,
+            low_cpu_mem_usage=low_cpu_mem_usage,
+            hotswap=hotswap,
+        )

    @classmethod
    # Copied from diffusers.loaders.lora_pipeline.SD3LoraLoaderMixin.load_lora_into_transformer with SD3Transformer2DModel->SkyReelsV2Transformer3DModel
@@ -6683,8 +6643,7 @@ class QwenImageLoraLoaderMixin(LoraBaseMixin):
            state_dict = {k: v for k, v in state_dict.items() if "dora_scale" not in k}

        has_alphas_in_sd = any(k.endswith(".alpha") for k in state_dict)
-        has_lora_unet = any(k.startswith("lora_unet_") for k in state_dict)
-        if has_alphas_in_sd or has_lora_unet:
+        if has_alphas_in_sd:
            state_dict = _convert_non_diffusers_qwen_lora_to_diffusers(state_dict)

        out = (state_dict, metadata) if return_lora_metadata else state_dict
@@ -299,7 +299,6 @@ class Decoder(nn.Module):
        act_fn: Union[str, Tuple[str]] = "silu",
        upsample_block_type: str = "pixel_shuffle",
        in_shortcut: bool = True,
-        conv_act_fn: str = "relu",
    ):
        super().__init__()

@@ -350,7 +349,7 @@ class Decoder(nn.Module):
        channels = block_out_channels[0] if layers_per_block[0] > 0 else block_out_channels[1]

        self.norm_out = RMSNorm(channels, 1e-5, elementwise_affine=True, bias=True)
-        self.conv_act = get_activation(conv_act_fn)
+        self.conv_act = nn.ReLU()
        self.conv_out = None

        if layers_per_block[0] > 0:
@@ -415,12 +414,6 @@ class AutoencoderDC(ModelMixin, ConfigMixin, FromOriginalModelMixin):
            The normalization type(s) to use in the decoder.
        decoder_act_fns (`Union[str, Tuple[str]]`, defaults to `"silu"`):
            The activation function(s) to use in the decoder.
-        encoder_out_shortcut  (`bool`, defaults to `True`):
-            Whether to use shortcut at the end of the encoder.
-        decoder_in_shortcut (`bool`, defaults to `True`):
-            Whether to use shortcut at the beginning of the decoder.
-        decoder_conv_act_fn (`str`, defaults to `"relu"`):
-            The activation function to use at the end of the decoder.
        scaling_factor (`float`, defaults to `1.0`):
            The multiplicative inverse of the root mean square of the latent features. This is used to scale the latent
            space to have unit variance when training the diffusion model. The latents are scaled with the formula `z =
@@ -448,9 +441,6 @@ class AutoencoderDC(ModelMixin, ConfigMixin, FromOriginalModelMixin):
        downsample_block_type: str = "pixel_unshuffle",
        decoder_norm_types: Union[str, Tuple[str]] = "rms_norm",
        decoder_act_fns: Union[str, Tuple[str]] = "silu",
-        encoder_out_shortcut: bool = True,
-        decoder_in_shortcut: bool = True,
-        decoder_conv_act_fn: str = "relu",
        scaling_factor: float = 1.0,
    ) -> None:
        super().__init__()
@@ -464,7 +454,6 @@ class AutoencoderDC(ModelMixin, ConfigMixin, FromOriginalModelMixin):
            layers_per_block=encoder_layers_per_block,
            qkv_multiscales=encoder_qkv_multiscales,
            downsample_block_type=downsample_block_type,
-            out_shortcut=encoder_out_shortcut,
        )
        self.decoder = Decoder(
            in_channels=in_channels,
@@ -477,8 +466,6 @@ class AutoencoderDC(ModelMixin, ConfigMixin, FromOriginalModelMixin):
            norm_type=decoder_norm_types,
            act_fn=decoder_act_fns,
            upsample_block_type=upsample_block_type,
-            in_shortcut=decoder_in_shortcut,
-            conv_act_fn=decoder_conv_act_fn,
        )

        self.spatial_compression_ratio = 2 ** (len(encoder_block_out_channels) - 1)
@@ -726,29 +726,23 @@ def _caching_allocator_warmup(
    very large margin.
    """
    factor = 2 if hf_quantizer is None else hf_quantizer.get_cuda_warm_up_factor()
-
-    # Keep only accelerator devices
+    # Remove disk and cpu devices, and cast to proper torch.device
    accelerator_device_map = {
        param: torch.device(device)
        for param, device in expanded_device_map.items()
        if str(device) not in ["cpu", "disk"]
    }
-    if not accelerator_device_map:
-        return
-
-    elements_per_device = defaultdict(int)
+    total_byte_count = defaultdict(lambda: 0)
    for param_name, device in accelerator_device_map.items():
        try:
-            p = model.get_parameter(param_name)
+            param = model.get_parameter(param_name)
        except AttributeError:
-            try:
-                p = model.get_buffer(param_name)
-            except AttributeError:
-                raise AttributeError(f"Parameter or buffer with name={param_name} not found in model")
+            param = model.get_buffer(param_name)
+        # The dtype of different parameters may be different with composite models or `keep_in_fp32_modules`
+        param_byte_count = param.numel() * param.element_size()
        # TODO: account for TP when needed.
-        elements_per_device[device] += p.numel()
+        total_byte_count[device] += param_byte_count

    # This will kick off the caching allocator to avoid having to Malloc afterwards
-    for device, elem_count in elements_per_device.items():
-        warmup_elems = max(1, elem_count // factor)
-        _ = torch.empty(warmup_elems, dtype=dtype, device=device, requires_grad=False)
+    for device, byte_count in total_byte_count.items():
+        _ = torch.empty(byte_count // factor, dtype=dtype, device=device, requires_grad=False)
@@ -290,7 +290,7 @@ class ModularPipelineBlocks(ConfigMixin, PushToHubMixin):
    def from_pretrained(
        cls,
        pretrained_model_name_or_path: str,
-        trust_remote_code: Optional[bool] = None,
+        trust_remote_code: bool = False,
        **kwargs,
    ):
        hub_kwargs_names = [
@@ -480,11 +480,6 @@ class QwenImagePipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
                of [Imagen Paper](https://huggingface.co/papers/2205.11487). Guidance scale is enabled by setting
                `guidance_scale > 1`. Higher guidance scale encourages to generate images that are closely linked to
                the text `prompt`, usually at the expense of lower image quality.
-
-                This parameter in the pipeline is there to support future guidance-distilled models when they come up.
-                Note that passing `guidance_scale` to the pipeline is ineffective. To enable classifier-free guidance,
-                please pass `true_cfg_scale` and `negative_prompt` (even an empty negative prompt like " ") should
-                enable classifier-free guidance computations.
            num_images_per_prompt (`int`, *optional*, defaults to 1):
                The number of images to generate per prompt.
            generator (`torch.Generator` or `List[torch.Generator]`, *optional*):
@@ -62,6 +62,25 @@ EXAMPLE_DOC_STRING = """
        >>> image.save("qwenimage_edit.png")
        ```
 """
+PREFERRED_QWENIMAGE_RESOLUTIONS = [
+    (672, 1568),
+    (688, 1504),
+    (720, 1456),
+    (752, 1392),
+    (800, 1328),
+    (832, 1248),
+    (880, 1184),
+    (944, 1104),
+    (1024, 1024),
+    (1104, 944),
+    (1184, 880),
+    (1248, 832),
+    (1328, 800),
+    (1392, 752),
+    (1456, 720),
+    (1504, 688),
+    (1568, 672),
+]


 # Copied from diffusers.pipelines.qwenimage.pipeline_qwenimage.calculate_shift
@@ -546,6 +565,7 @@ class QwenImageEditPipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
        callback_on_step_end: Optional[Callable[[int, int, Dict], None]] = None,
        callback_on_step_end_tensor_inputs: List[str] = ["latents"],
        max_sequence_length: int = 512,
+        _auto_resize: bool = True,
    ):
        r"""
        Function invoked when calling the pipeline for generation.
@@ -577,11 +597,6 @@ class QwenImageEditPipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
                of [Imagen Paper](https://huggingface.co/papers/2205.11487). Guidance scale is enabled by setting
                `guidance_scale > 1`. Higher guidance scale encourages to generate images that are closely linked to
                the text `prompt`, usually at the expense of lower image quality.
-
-                This parameter in the pipeline is there to support future guidance-distilled models when they come up.
-                Note that passing `guidance_scale` to the pipeline is ineffective. To enable classifier-free guidance,
-                please pass `true_cfg_scale` and `negative_prompt` (even an empty negative prompt like " ") should
-                enable classifier-free guidance computations.
            num_images_per_prompt (`int`, *optional*, defaults to 1):
                The number of images to generate per prompt.
            generator (`torch.Generator` or `List[torch.Generator]`, *optional*):
@@ -626,7 +641,8 @@ class QwenImageEditPipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
            returning a tuple, the first element is a list with the generated images.
        """
        image_size = image[0].size if isinstance(image, list) else image.size
-        calculated_width, calculated_height, _ = calculate_dimensions(1024 * 1024, image_size[0] / image_size[1])
+        width, height = image_size
+        calculated_width, calculated_height, _ = calculate_dimensions(1024 * 1024, width / height)
        height = height or calculated_height
        width = width or calculated_width

@@ -664,9 +680,18 @@ class QwenImageEditPipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
        device = self._execution_device
        # 3. Preprocess image
        if image is not None and not (isinstance(image, torch.Tensor) and image.size(1) == self.latent_channels):
-            image = self.image_processor.resize(image, calculated_height, calculated_width)
+            img = image[0] if isinstance(image, list) else image
+            image_height, image_width = self.image_processor.get_default_height_width(img)
+            aspect_ratio = image_width / image_height
+            if _auto_resize:
+                _, image_width, image_height = min(
+                    (abs(aspect_ratio - w / h), w, h) for w, h in PREFERRED_QWENIMAGE_RESOLUTIONS
+                )
+            image_width = image_width // multiple_of * multiple_of
+            image_height = image_height // multiple_of * multiple_of
+            image = self.image_processor.resize(image, image_height, image_width)
            prompt_image = image
-            image = self.image_processor.preprocess(image, calculated_height, calculated_width)
+            image = self.image_processor.preprocess(image, image_height, image_width)
            image = image.unsqueeze(2)

        has_neg_prompt = negative_prompt is not None or (
@@ -683,6 +708,9 @@ class QwenImageEditPipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
            max_sequence_length=max_sequence_length,
        )
        if do_true_cfg:
+            # negative image is the same size as the original image, but all pixels are white
+            # negative_image = Image.new("RGB", (image.width, image.height), (255, 255, 255))
+
            negative_prompt_embeds, negative_prompt_embeds_mask = self.encode_prompt(
                image=prompt_image,
                prompt=negative_prompt,
@@ -709,7 +737,7 @@ class QwenImageEditPipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
        img_shapes = [
            [
                (1, height // self.vae_scale_factor // 2, width // self.vae_scale_factor // 2),
-                (1, calculated_height // self.vae_scale_factor // 2, calculated_width // self.vae_scale_factor // 2),
+                (1, image_height // self.vae_scale_factor // 2, image_width // self.vae_scale_factor // 2),
            ]
        ] * batch_size

@@ -568,11 +568,6 @@ class QwenImageImg2ImgPipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
                of [Imagen Paper](https://huggingface.co/papers/2205.11487). Guidance scale is enabled by setting
                `guidance_scale > 1`. Higher guidance scale encourages to generate images that are closely linked to
                the text `prompt`, usually at the expense of lower image quality.
-
-                This parameter in the pipeline is there to support future guidance-distilled models when they come up.
-                Note that passing `guidance_scale` to the pipeline is ineffective. To enable classifier-free guidance,
-                please pass `true_cfg_scale` and `negative_prompt` (even an empty negative prompt like " ") should
-                enable classifier-free guidance computations.
            num_images_per_prompt (`int`, *optional*, defaults to 1):
                The number of images to generate per prompt.
            generator (`torch.Generator` or `List[torch.Generator]`, *optional*):
@@ -698,11 +698,6 @@ class QwenImageInpaintPipeline(DiffusionPipeline, QwenImageLoraLoaderMixin):
                of [Imagen Paper](https://huggingface.co/papers/2205.11487). Guidance scale is enabled by setting
                `guidance_scale > 1`. Higher guidance scale encourages to generate images that are closely linked to
                the text `prompt`, usually at the expense of lower image quality.
-
-                This parameter in the pipeline is there to support future guidance-distilled models when they come up.
-                Note that passing `guidance_scale` to the pipeline is ineffective. To enable classifier-free guidance,
-                please pass `true_cfg_scale` and `negative_prompt` (even an empty negative prompt like " ") should
-                enable classifier-free guidance computations.
            num_images_per_prompt (`int`, *optional*, defaults to 1):
                The number of images to generate per prompt.
            generator (`torch.Generator` or `List[torch.Generator]`, *optional*):
@@ -339,8 +339,7 @@ def offload_models(
            original_devices = [next(m.parameters()).device for m in modules]
        else:
            assert len(modules) == 1
-            # For DiffusionPipeline, wrap the device in a list to make it iterable
-            original_devices = [modules[0].device]
+            original_devices = modules[0].device
        # move to target device
        for m in modules:
            m.to(device)
@@ -45,6 +45,7 @@ DIFFUSERS_ATTN_BACKEND = os.getenv("DIFFUSERS_ATTN_BACKEND", "native")
 DIFFUSERS_ATTN_CHECKS = os.getenv("DIFFUSERS_ATTN_CHECKS", "0") in ENV_VARS_TRUE_VALUES
 DEFAULT_HF_PARALLEL_LOADING_WORKERS = 8
 HF_ENABLE_PARALLEL_LOADING = os.environ.get("HF_ENABLE_PARALLEL_LOADING", "").upper() in ENV_VARS_TRUE_VALUES
+DIFFUSERS_DISABLE_REMOTE_CODE = os.getenv("DIFFUSERS_DISABLE_REMOTE_CODE", "false").lower() in ENV_VARS_TRUE_VALUES

 # Below should be `True` if the current version of `peft` and `transformers` are compatible with
 # PEFT backend. Will automatically fall back to PEFT backend if the correct versions of the libraries are
@@ -20,7 +20,6 @@ import json
 import os
 import re
 import shutil
-import signal
 import sys
 import threading
 from pathlib import Path
@@ -34,6 +33,7 @@ from packaging import version

 from .. import __version__
 from . import DIFFUSERS_DYNAMIC_MODULE_NAME, HF_MODULES_CACHE, logging
+from .constants import DIFFUSERS_DISABLE_REMOTE_CODE


 logger = logging.get_logger(__name__)  # pylint: disable=invalid-name
@@ -159,52 +159,25 @@ def check_imports(filename):
    return get_relative_imports(filename)


-def _raise_timeout_error(signum, frame):
-    raise ValueError(
-        "Loading this model requires you to execute custom code contained in the model repository on your local "
-        "machine. Please set the option `trust_remote_code=True` to permit loading of this model."
-    )
-
-
 def resolve_trust_remote_code(trust_remote_code, model_name, has_remote_code):
-    if trust_remote_code is None:
-        if has_remote_code and TIME_OUT_REMOTE_CODE > 0:
-            prev_sig_handler = None
-            try:
-                prev_sig_handler = signal.signal(signal.SIGALRM, _raise_timeout_error)
-                signal.alarm(TIME_OUT_REMOTE_CODE)
-                while trust_remote_code is None:
-                    answer = input(
-                        f"The repository for {model_name} contains custom code which must be executed to correctly "
-                        f"load the model. You can inspect the repository content at https://hf.co/{model_name}.\n"
-                        f"You can avoid this prompt in future by passing the argument `trust_remote_code=True`.\n\n"
-                        f"Do you wish to run the custom code? [y/N] "
-                    )
-                    if answer.lower() in ["yes", "y", "1"]:
-                        trust_remote_code = True
-                    elif answer.lower() in ["no", "n", "0", ""]:
-                        trust_remote_code = False
-                signal.alarm(0)
-            except Exception:
-                # OS which does not support signal.SIGALRM
-                raise ValueError(
-                    f"The repository for {model_name} contains custom code which must be executed to correctly "
-                    f"load the model. You can inspect the repository content at https://hf.co/{model_name}.\n"
-                    f"Please pass the argument `trust_remote_code=True` to allow custom code to be run."
-                )
-            finally:
-                if prev_sig_handler is not None:
-                    signal.signal(signal.SIGALRM, prev_sig_handler)
-                    signal.alarm(0)
-        elif has_remote_code:
-            # For the CI which puts the timeout at 0
-            _raise_timeout_error(None, None)
+    trust_remote_code = trust_remote_code and not DIFFUSERS_DISABLE_REMOTE_CODE
+    if DIFFUSERS_DISABLE_REMOTE_CODE:
+        logger.warning(
+            "Downloading remote code is disabled globally via the DIFFUSERS_DISABLE_REMOTE_CODE environment variable. Ignoring `trust_remote_code`."
+        )

    if has_remote_code and not trust_remote_code:
-        raise ValueError(
-            f"Loading {model_name} requires you to execute the configuration file in that"
-            " repo on your local machine. Make sure you have read the code there to avoid malicious use, then"
-            " set the option `trust_remote_code=True` to remove this error."
+        error_msg = f"The repository for {model_name} contains custom code. "
+        error_msg += (
+            "Downloading remote code is disabled globally via the DIFFUSERS_DISABLE_REMOTE_CODE environment variable."
+            if DIFFUSERS_DISABLE_REMOTE_CODE
+            else "Pass `trust_remote_code=True` to allow loading remote code modules."
+        )
+        raise ValueError(error_msg)
+
+    elif has_remote_code and trust_remote_code:
+        logger.warning(
+            f"`trust_remote_code` is enabled. Downloading code from {model_name}. Please ensure you trust the contents of this repository"
        )

    return trust_remote_code
Author	SHA1	Message	Date
DN6	50c7ddeaea	update	2025-08-18 14:02:05 +05:30
DN6	4b2b2b221b	update	2025-08-18 13:29:07 +05:30
DN6	0db2ea2bc8	update	2025-08-18 13:20:59 +05:30
DN6	1b26e309f4	update	2025-08-18 11:40:02 +05:30