Merge 6fd6ffd023 into cd8c7a2306

Throttle dynamic VRAM prepare logging (#13704 )
Revert "Fix Content-Disposition header missing 'attachment;' prefix (#13093 )" (#13733 )
2026-05-23 15:37:27 +08:00 · 2026-05-07 10:41:24 +08:00 · 2026-05-07 10:41:13 +08:00 · 2026-05-06 10:08:35 -07:00 · 2025-04-10 18:16:20 +02:00 · 2025-04-08 19:27:56 +02:00
4 changed files with 13 additions and 5 deletions
--- a/comfy/cli_args.py
+++ b/comfy/cli_args.py
@ -157,6 +157,7 @@ parser.add_argument("--force-non-blocking", action="store_true", help="Force Com

 parser.add_argument("--default-hashing-function", type=str, choices=['md5', 'sha1', 'sha256', 'sha512'], default='sha256', help="Allows you to choose the hash function to use for duplicate filename / contents comparison. Default is sha256.")

+parser.add_argument("--disable-fp8-compute", action="store_true", help="Prevent ComfyUI from activating fp8 compute in Nvidia cards that support it. Can prevent some issues with some models not suitable for fp8 compute.")
 parser.add_argument("--disable-smart-memory", action="store_true", help="Force ComfyUI to agressively offload to regular ram instead of keeping models in vram when it can.")
 parser.add_argument("--deterministic", action="store_true", help="Make pytorch use slower deterministic algorithms when it can. Note that this might not make images deterministic in all cases.")

--- a/comfy/model_management.py
+++ b/comfy/model_management.py
@ -1699,6 +1699,8 @@ def supports_fp8_compute(device=None):

    if not is_nvidia():
        return False
+    if args.disable_fp8_compute:
+        return False

    props = torch.cuda.get_device_properties(device)
    if props.major >= 9:
--- a/comfy/model_patcher.py
+++ b/comfy/model_patcher.py
@ -26,6 +26,7 @@ import uuid
 from typing import Callable, Optional

 import torch
+import tqdm

 import comfy.float
 import comfy.hooks
@ -1651,7 +1652,11 @@ class ModelPatcherDynamic(ModelPatcher):
                self.model.model_loaded_weight_memory += casted_buf.numel() * casted_buf.element_size()

            force_load_stat = f" Force pre-loaded {len(self.backup)} weights: {self.model.model_loaded_weight_memory // 1024} KB." if len(self.backup) > 0 else ""
-            logging.info(f"Model {self.model.__class__.__name__} prepared for dynamic VRAM loading. {allocated_size // (1024 ** 2)}MB Staged. {num_patches} patches attached.{force_load_stat}")
+            log_key = (self.patches_uuid, allocated_size, num_patches, len(self.backup), self.model.model_loaded_weight_memory)
+            in_loop = bool(getattr(tqdm.tqdm, "_instances", None))
+            level = logging.DEBUG if in_loop and getattr(self, "_last_prepare_log_key", None) == log_key else logging.INFO
+            self._last_prepare_log_key = log_key
+            logging.log(level, f"Model {self.model.__class__.__name__} prepared for dynamic VRAM loading. {allocated_size // (1024 ** 2)}MB Staged. {num_patches} patches attached.{force_load_stat}")

            self.model.device = device_to
            self.model.current_weight_patches_uuid = self.patches_uuid
--- a/server.py
+++ b/server.py
@ -560,7 +560,7 @@ class PromptServer():
                            buffer.seek(0)

                            return web.Response(body=buffer.read(), content_type=f'image/{image_format}',
-                                                headers={"Content-Disposition": f"attachment; filename=\"{filename}\""})
+                                                headers={"Content-Disposition": f"filename=\"{filename}\""})

                    if 'channel' not in request.rel_url.query:
                        channel = 'rgba'
@ -580,7 +580,7 @@ class PromptServer():
                            buffer.seek(0)

                            return web.Response(body=buffer.read(), content_type='image/png',
-                                                headers={"Content-Disposition": f"attachment; filename=\"{filename}\""})
+                                                headers={"Content-Disposition": f"filename=\"{filename}\""})

                    elif channel == 'a':
                        with Image.open(file) as img:
@ -597,7 +597,7 @@ class PromptServer():
                            alpha_buffer.seek(0)

                            return web.Response(body=alpha_buffer.read(), content_type='image/png',
-                                                headers={"Content-Disposition": f"attachment; filename=\"{filename}\""})
+                                                headers={"Content-Disposition": f"filename=\"{filename}\""})
                    else:
                        # Use the content type from asset resolution if available,
                        # otherwise guess from the filename.
@ -614,7 +614,7 @@ class PromptServer():
                        return web.FileResponse(
                            file,
                            headers={
-                                "Content-Disposition": f"attachment; filename=\"{filename}\"",
+                                "Content-Disposition": f"filename=\"{filename}\"",
                                "Content-Type": content_type
                            }
                        )
Author	SHA1	Message	Date
Silver	cfc159022e	Merge `6fd6ffd023` into `cd8c7a2306`	2026-05-07 10:41:24 +08:00
Jukka Seppänen	cd8c7a2306	Throttle dynamic VRAM prepare logging (#13704 )	2026-05-07 10:41:13 +08:00
guill	6bcd8b96ab	Revert "Fix Content-Disposition header missing 'attachment;' prefix (#13093 )" (#13733 ) Some checks are pending Python Linting / Run Ruff (push) Waiting to run Details Python Linting / Run Pylint (push) Waiting to run Details Full Comfy CI Workflow Runs / test-stable (12.1, , linux, 3.10, [self-hosted Linux], stable) (push) Waiting to run Details Full Comfy CI Workflow Runs / test-stable (12.1, , linux, 3.11, [self-hosted Linux], stable) (push) Waiting to run Details Full Comfy CI Workflow Runs / test-stable (12.1, , linux, 3.12, [self-hosted Linux], stable) (push) Waiting to run Details Full Comfy CI Workflow Runs / test-unix-nightly (12.1, , linux, 3.11, [self-hosted Linux], nightly) (push) Waiting to run Details Execution Tests / test (macos-latest) (push) Waiting to run Details Execution Tests / test (ubuntu-latest) (push) Waiting to run Details Execution Tests / test (windows-latest) (push) Waiting to run Details Test server launches without errors / test (push) Waiting to run Details Unit Tests / test (macos-latest) (push) Waiting to run Details Unit Tests / test (ubuntu-latest) (push) Waiting to run Details Unit Tests / test (windows-2022) (push) Waiting to run Details This reverts commit `ea6880b04b`.	2026-05-06 10:08:35 -07:00
Silver	6fd6ffd023	Merge branch 'comfyanonymous:master' into fp8compute_disable	2025-04-10 18:16:20 +02:00
silveroxides	a6b22bd779	Add launch argument for disabling fp8 compute	2025-04-08 19:27:56 +02:00