# Copyright 2024 Bytedance Ltd. and/or its affiliates # # Licensed under the Apache License, Version 2.0 (the "License"); # you may not use this file except in compliance with the License. # You may obtain a copy of the License at # # http://www.apache.org/licenses/LICENSE-2.0 # # Unless required by applicable law or agreed to in writing, software # distributed under the License is distributed on an "AS IS" BASIS, # WITHOUT WARRANTIES OR CONDITIONS OF ANY KIND, either express or implied. # See the License for the specific language governing permissions and # limitations under the License. import types import pytest import torch from verl.utils.model import extract_multi_modal_inputs from verl.utils.tokenizer import build_multimodal_processor_inputs def test_build_messages_replaces_audio_placeholder() -> None: pytest.importorskip("datasets") from verl.utils.dataset.rl_dataset import RLHFDataset dataset = RLHFDataset.__new__(RLHFDataset) dataset.prompt_key = "prompt" dataset.image_key = "images" dataset.video_key = "videos" dataset.audio_key = "audios" dataset.processor = object() example = { "prompt": [ {"role": "user", "content": "Listen to this: