Offloaded cache: fix generate (#34921)

* fix cache impl * require_torch_gpu * fix mamba * fix copies
2024-11-28 15:05:56 +01:00
parent 57ca9e6d2f
commit 5e8c1d713d
6 changed files with 91 additions and 19 deletions
--- a/tests/models/whisper/test_modeling_whisper.py
+++ b/tests/models/whisper/test_modeling_whisper.py
@@ -567,6 +567,12 @@ class WhisperModelTest(ModelTesterMixin, GenerationTesterMixin, PipelineTesterMi
    def test_generate_with_head_masking(self):
        pass

+    @parameterized.expand([("offloaded",)])
+    @pytest.mark.generate
+    @unittest.skip(reason="Whisper doesnt work with offloaded cache implementation yet")
+    def test_offloaded_cache_implementation(self, cache_implementation):
+        pass
+
    @require_torch_fp16
    def test_generate_fp16(self):
        config, input_dict = self.model_tester.prepare_config_and_inputs()