@misc{indiciae7e0865a3ca95, title = {A Spark of Vision-Language Intelligence: 2-Dimensional Autoregressive Transformer for Efficient Finegrained Image Generation}, author = {Liang Chen and Sinan Tan and Zefan Cai and Weichu Xie and Haozhe Zhao and Yichi Zhang and Junyang Lin and Jinze Bai and Tianyu Liu and Baobao Chang}, year = {2024}, url = {https://arxiv.org/abs/2410.01912}, note = {Source identifier: 2410.01912} }