fix: use the base layers weight in mistral rocm (#2155)

b966bc0d · drbh · GitHub · 5d97e0c4 · b966bc0d
Unverified Commit b966bc0d authored Jul 02, 2024 by drbh Committed by GitHub Jul 02, 2024
Hide whitespace changes
Inline Side-by-side

Showing with 3 additions and 1 deletion

server/text_generation_server/models/custom_modeling/flash_mistral_modeling.py ...n_server/models/custom_modeling/flash_mistral_modeling.py +3 -1

No files found.
--- a/server/text_generation_server/models/custom_modeling/flash_mistral_modeling.py
+++ b/server/text_generation_server/models/custom_modeling/flash_mistral_modeling.py
@@ -315,7 +315,9 @@ class MistralMLP(nn.Module):
                dtype=hidden_states.dtype,
                device="cuda",
            )
-            _custom_C.LLMM_Silu(self.gate_up_proj.linear.weight, hidden_states, out, 8)
+            _custom_C.LLMM_Silu(
+                self.gate_up_proj.base_layer.linear.weight, hidden_states, out, 8
+            )
            return self.down_proj(out, adapter_data)
        else:
            gate_up_states = self.gate_up_proj(hidden_states, adapter_data)