@misc{indiciae0641ea65267f, title = {How and What to Imagine? Visual Thinking in Unified Multimodal Models for Cross-View Spatial Reasoning}, author = {Qian Yang and Ankur Sikarwar and Huy Le and Le Zhang and Zhuan Shi and Perouz Taslakian and Aishwarya Agrawal}, year = {2026}, url = {https://arxiv.org/abs/2605.27310}, note = {Source identifier: 2605.27310} }