@misc{indiciaee78f74a203ad, title = {SketchVLM: Vision language models can annotate images to explain thoughts and guide users}, author = {Brandon Collins and Logan Bolton and Hung Huy Nguyen and Mohammad Reza Taesiri and Trung Bui and Anh Totti Nguyen}, year = {2026}, url = {https://arxiv.org/abs/2604.22875}, note = {Source identifier: 2604.22875} }