@misc{indiciae80d64e25f994, title = {Omni-o3: Deep Nested Omnimodal Deduction for Deliberative Audio-Visual Reasoning}, author = {Zhicheng Zhang and Wentao Gu and Weicheng Wang and Yongjie Zhu and Wenyu Qin and Meng Wang and Pengfei Wan and Jufeng Yang}, year = {2026}, url = {https://arxiv.org/abs/2604.24191}, note = {Source identifier: 2604.24191} }