@misc{indiciae1c315d404f14, title = {An Efficient vLLM-Based Inference Pipeline for Unified Audio Understanding and Generation}, author = {Haoran Wang and Jinchuan Tian and Siddhant Arora and Shinji Watanabe}, year = {2026}, url = {https://arxiv.org/abs/2607.02119}, note = {Source identifier: 2607.02119} }