Uploaded sanitized BC250/ROCm Repository.
This commit is contained in:
@@ -0,0 +1,186 @@
|
||||
"""Remove --cpu-vae: Let VAE run on GPU (only 320MB, easily fits). Restart ComfyUI and test."""
|
||||
import paramiko, time, json, textwrap
|
||||
|
||||
k = paramiko.Ed25519Key.from_private_key_file(r'C:\Users\fabia\.ssh\id_ed25519')
|
||||
c = paramiko.SSHClient()
|
||||
c.set_missing_host_key_policy(paramiko.AutoAddPolicy())
|
||||
c.connect('192.168.178.150', username='fabian', pkey=k, timeout=15)
|
||||
|
||||
def sh(cmd, timeout=60):
|
||||
chan = c.get_transport().open_session()
|
||||
chan.settimeout(timeout)
|
||||
chan.exec_command('/bin/bash -l -c ' + "'" + cmd.replace("'", "'\\''") + "'")
|
||||
out = b""
|
||||
while True:
|
||||
try:
|
||||
chunk = chan.recv(65536)
|
||||
if not chunk: break
|
||||
out += chunk
|
||||
except: break
|
||||
chan.close()
|
||||
return out.decode(errors='replace').strip()
|
||||
|
||||
def sftp_write(path, content):
|
||||
sftp = c.open_sftp()
|
||||
with sftp.open(path, 'w') as f:
|
||||
f.write(content)
|
||||
sftp.close()
|
||||
|
||||
def sftp_read(path):
|
||||
sftp = c.open_sftp()
|
||||
with sftp.open(path, 'r') as f:
|
||||
data = f.read().decode(errors='replace')
|
||||
sftp.close()
|
||||
return data
|
||||
|
||||
# Kill old
|
||||
print("Killing ComfyUI...")
|
||||
sh('pkill -9 -f "python3.*main.py" 2>/dev/null; sleep 2')
|
||||
|
||||
# New launcher WITHOUT --cpu-vae
|
||||
launcher = textwrap.dedent("""\
|
||||
#!/bin/bash
|
||||
export HSA_OVERRIDE_GFX_VERSION=10.1.0
|
||||
export HIP_VISIBLE_DEVICES=0
|
||||
export HSA_ENABLE_SDMA=0
|
||||
export HSA_TOOLS_LIB=""
|
||||
export HSA_TOOLS_REPORT_LOAD_FAILURE=0
|
||||
export PYTORCH_HIP_ALLOC_CONF=expandable_segments:False
|
||||
export OMP_NUM_THREADS=12
|
||||
export MKL_NUM_THREADS=12
|
||||
export OPENBLAS_NUM_THREADS=12
|
||||
export MIOPEN_FIND_MODE=1
|
||||
|
||||
cd ~/ComfyUI
|
||||
source ~/comfyui-env/bin/activate
|
||||
|
||||
# --novram: model weights on CPU, GPU computes (correct for shared-memory APU)
|
||||
# --force-fp16: half precision
|
||||
# NO --cpu-vae: VAE is only 320MB, runs fine on GPU and much faster
|
||||
exec python3 main.py \\
|
||||
--listen 0.0.0.0 --port 8188 \\
|
||||
--novram \\
|
||||
--force-fp16 \\
|
||||
--disable-smart-memory
|
||||
""")
|
||||
sftp_write('/tmp/run_comfyui.sh', launcher)
|
||||
sh('chmod +x /tmp/run_comfyui.sh')
|
||||
print("Launcher updated: --novram --force-fp16 (NO --cpu-vae)")
|
||||
|
||||
# Start
|
||||
sh('rm -f /tmp/comfyui.log; touch /tmp/comfyui.log')
|
||||
sh('nohup /tmp/run_comfyui.sh > /tmp/comfyui.log 2>&1 &')
|
||||
time.sleep(3)
|
||||
pid = sh('pgrep -f "python3.*main.py"')
|
||||
print(f"PID: {pid}")
|
||||
|
||||
# Wait for ready
|
||||
print("Waiting for server...", end='', flush=True)
|
||||
for i in range(120):
|
||||
code = sh('curl -s -o /dev/null -w "%{http_code}" http://127.0.0.1:8188/ 2>/dev/null', timeout=5)
|
||||
if '200' in code:
|
||||
print(f" READY ({i*2}s)")
|
||||
break
|
||||
if i % 10 == 0 and i > 0:
|
||||
log = sftp_read('/tmp/comfyui.log')
|
||||
lines = [l for l in log.split('\n') if l.strip()]
|
||||
print(f"\n [{i*2}s] {lines[-1][:80] if lines else '...'}", end='', flush=True)
|
||||
else:
|
||||
print('.', end='', flush=True)
|
||||
time.sleep(2)
|
||||
|
||||
# Submit workflow
|
||||
print("\nSubmitting workflow...")
|
||||
workflow = {
|
||||
"prompt": {
|
||||
"1": {"class_type": "UnetLoaderGGUF", "inputs": {"unet_name": "z_image_turbo-Q5_K_S.gguf"}},
|
||||
"2": {"class_type": "CLIPLoaderGGUF", "inputs": {"clip_name": "Qwen3-4B.i1-Q5_K_S.gguf", "type": "qwen_image"}},
|
||||
"3": {"class_type": "VAELoader", "inputs": {"vae_name": "ae.safetensors"}},
|
||||
"4": {"class_type": "CLIPTextEncode", "inputs": {"text": "A red fox in a snowy forest, photorealistic", "clip": ["2", 0]}},
|
||||
"9": {"class_type": "CLIPTextEncode", "inputs": {"text": "", "clip": ["2", 0]}},
|
||||
"5": {"class_type": "EmptyLatentImage", "inputs": {"width": 512, "height": 512, "batch_size": 1}},
|
||||
"6": {"class_type": "KSampler", "inputs": {
|
||||
"model": ["1", 0], "positive": ["4", 0], "negative": ["9", 0],
|
||||
"latent_image": ["5", 0], "seed": 42, "steps": 8, "cfg": 1.0,
|
||||
"sampler_name": "euler", "scheduler": "simple", "denoise": 1.0
|
||||
}},
|
||||
"7": {"class_type": "VAEDecode", "inputs": {"samples": ["6", 0], "vae": ["3", 0]}},
|
||||
"8": {"class_type": "SaveImage", "inputs": {"images": ["7", 0], "filename_prefix": "ZImageTurbo_v2"}}
|
||||
}
|
||||
}
|
||||
sftp_write('/tmp/wf.json', json.dumps(workflow))
|
||||
resp = sh('curl -s -X POST http://127.0.0.1:8188/prompt -H "Content-Type: application/json" -d @/tmp/wf.json', timeout=10)
|
||||
print(f"Response: {resp[:150]}")
|
||||
|
||||
# Monitor
|
||||
print("\nMonitoring...")
|
||||
t0 = time.time()
|
||||
for i in range(120):
|
||||
elapsed = int(time.time() - t0)
|
||||
|
||||
gpu_temp = sh('cat /sys/class/drm/card0/device/hwmon/hwmon*/temp1_input 2>/dev/null', timeout=5)
|
||||
temp_c = int(gpu_temp) // 1000 if gpu_temp.isdigit() else '?'
|
||||
|
||||
try:
|
||||
log = sftp_read('/tmp/comfyui.log')
|
||||
except:
|
||||
log = ''
|
||||
|
||||
# Find last meaningful line
|
||||
last = ''
|
||||
sampling = ''
|
||||
for line in log.split('\n'):
|
||||
s = line.strip()
|
||||
if '/8' in s and ('it/s' in s or 's/it' in s):
|
||||
sampling = s
|
||||
if s and 'FETCH' not in s and 'startup tasks' not in s and 'DEPRECATION' not in s:
|
||||
last = s
|
||||
|
||||
display = sampling if sampling else last[-100:]
|
||||
print(f" [{elapsed:>4}s] {temp_c}C | {display}")
|
||||
|
||||
# Check output
|
||||
imgs = sh('ls ~/ComfyUI/output/ZImageTurbo_v2*.png 2>/dev/null', timeout=5)
|
||||
if imgs:
|
||||
print(f"\n*** IMAGE GENERATED! ***")
|
||||
print(f"File: {imgs}")
|
||||
print(f"Total time: {elapsed}s")
|
||||
# Show timing from log
|
||||
for line in log.split('\n')[-15:]:
|
||||
s = line.strip()
|
||||
if s and 'FETCH' not in s and 'startup' not in s:
|
||||
print(f" {s}")
|
||||
break
|
||||
|
||||
# Queue check
|
||||
q = sh('curl -s http://127.0.0.1:8188/queue 2>/dev/null', timeout=5)
|
||||
try:
|
||||
qd = json.loads(q)
|
||||
if not qd.get('queue_running') and not qd.get('queue_pending') and elapsed > 30:
|
||||
time.sleep(3)
|
||||
imgs = sh('ls ~/ComfyUI/output/ZImageTurbo_v2*.png 2>/dev/null', timeout=5)
|
||||
if imgs:
|
||||
print(f"\n*** IMAGE GENERATED! ***")
|
||||
print(f"File: {imgs}")
|
||||
print(f"Total time: {elapsed}s")
|
||||
else:
|
||||
print(f"\nQueue empty, no image:")
|
||||
for line in log.split('\n')[-20:]:
|
||||
if line.strip():
|
||||
print(f" {line.strip()}")
|
||||
break
|
||||
except:
|
||||
pass
|
||||
|
||||
alive = sh('pgrep -f "python3.*main.py" >/dev/null && echo Y || echo N', timeout=5)
|
||||
if alive == 'N':
|
||||
print("\n*** CRASHED ***")
|
||||
for line in log.split('\n')[-30:]:
|
||||
if line.strip():
|
||||
print(f" {line.strip()}")
|
||||
break
|
||||
|
||||
time.sleep(15)
|
||||
|
||||
c.close()
|
||||
print("\nDone.")
|
||||
Reference in New Issue
Block a user