mirror of
https://github.com/wassname/vllm.git
synced 2026-09-30 11:45:09 +08:00
[Bugfix][TPU] Fix TPU sampler output (#5978)
This commit is contained in:
1 parent
7041de4384
commit
54814fd85b
1 file changed
+1
-1
@@ -215,7 +215,7 @@ class TPUWorker(LoraNotSupportedWorkerBase):
|
||||
assert len(seq_group_metadata_list) > 0
|
||||
output = self.model_runner.execute_model(seq_group_metadata_list,
|
||||
self.tpu_cache)
|
||||
return [output]
|
||||
return output
|
||||
|
||||
def cache_swap(
|
||||
self,
|
||||
|
||||
Reference in new issue
Block a user