@misc{indiciae273c3cfd40b7, title = {Image-of-Thought Prompting for Visual Reasoning Refinement in Multimodal Large Language Models}, author = {Qiji Zhou and Ruochen Zhou and Zike Hu and Panzhong Lu and Siyang Gao and Yue Zhang}, year = {2024}, url = {https://arxiv.org/abs/2405.13872}, note = {Source identifier: 2405.13872} }