@misc{indiciae24641743dd8f, title = {High-Fidelity Text-to-Image Generation from Pre-Trained Vision-Language Models via Distribution-Conditioned Diffusion Decoding}, author = {Ji Woo Hong and Hee Suk Yoon and Gwanhyeong Koo and Eunseop Yoon and SooHwan Eom and Qi Dai and Chong Luo and Chang D. Yoo}, year = {2026}, url = {https://arxiv.org/abs/2603.13389}, note = {Source identifier: 2603.13389} }