@misc{indiciae2c7691c55c97, title = {Mind-VLA: Instruction-Aware Spatial Representation Alignment for Vision-Language-Action Models}, author = {Xingyu Ding and Yuzhong Zhao and Yang Wu and Chunhai Zhao and Chaoyang Zhao and Yifan Zhang and Jian Cheng}, year = {2026}, url = {https://arxiv.org/abs/2608.04633}, note = {Source identifier: 2608.04633} }