attn mask bugfix

This commit is contained in:
Kijai
2024-08-21 16:08:46 +03:00
parent c715144f4d
commit 617c33da0f
+3 -2
View File
@@ -707,9 +707,10 @@ class DoubleStreamBlock(nn.Module):
# make attention mask if not None
attn_mask = None
if txt_attention_mask is not None:
attn_mask = txt_attention_mask # b, seq_len
# F.scaled_dot_product_attention expects attn_mask to be bool for binary mask
attn_mask = txt_attention_mask.to(torch.bool) # b, seq_len
attn_mask = torch.cat(
(attn_mask, torch.ones(attn_mask.shape[0], img.shape[1]).to(attn_mask.device)), dim=1
(attn_mask, torch.ones(attn_mask.shape[0], img.shape[1], device=attn_mask.device, dtype=torch.bool)), dim=1
) # b, seq_len + img_len
# broadcast attn_mask to all heads