Speed up generation time a tiny bit.

This commit is contained in:
comfyanonymous
2024-06-18 10:22:53 -04:00
parent f0c52f4460
commit a9a6923922
2 changed files with 2 additions and 2 deletions
+1 -1
View File
@@ -1,7 +1,7 @@
[project]
name = "comfyui_tensorrt"
description = "TensorRT Node for ComfyUI\nThis node enables the best performance on NVIDIA RTX™ Graphics Cards (GPUs) for Stable Diffusion by leveraging NVIDIA TensorRT."
version = "0.1.2"
version = "0.1.3"
license = "LICENSE"
dependencies = [
"tensorrt>=10.0.1",
+1 -1
View File
@@ -96,7 +96,7 @@ class TrTUnet:
x = model_inputs_converted[k]
self.context.set_tensor_address(k, x[(x.shape[0] // curr_split_batch) * i:].data_ptr())
self.context.execute_async_v3(stream_handle=stream.cuda_stream)
stream.synchronize()
# stream.synchronize() #don't need to sync stream since it's the default torch one
return out
def load_state_dict(self, sd, strict=False):