@misc{indiciae19d43b23ffa1, title = {Metropolis-Hastings Captioning Game: Knowledge Fusion of Vision Language Models via Decentralized Bayesian Inference}, author = {Yuta Matsui and Ryosuke Yamaki and Ryo Ueda and Seitaro Shinagawa and Tadahiro Taniguchi}, year = {2025}, url = {https://arxiv.org/abs/2504.09620}, note = {Source identifier: 2504.09620} }