docker run --rm --gpus all -p 8000:8000 \ -v /path/to/models:/models \ vllm/vllm-openai:latest \ --model /models/my-fine-tuned-llama __ __