[Bugfix] Standardize merging multimodal embeddings (#26771)

Signed-off-by: DarkLight1337 <tlleungac@connect.ust.hk>
This commit is contained in:
Cyrus Leung
2025-10-14 17:36:21 +08:00
committed by GitHub
parent 577d498212
commit d2f816d6ff
19 changed files with 57 additions and 57 deletions

View File

@@ -762,7 +762,7 @@ class MiniCPMO(MiniCPMV2_6):
for modality in modalities:
if modality == "audios":
audio_input = modalities["audios"]
audio_features = self._process_audio_input(audio_input)
multimodal_embeddings += tuple(audio_features)
audio_embeddings = self._process_audio_input(audio_input)
multimodal_embeddings += tuple(audio_embeddings)
return multimodal_embeddings