docker run --rm --gpus all -p 8000:8000 \ vllm/vllm-openai:latest \ --model mistralai/Mistral-7B-v0.1 \ --max-model-len 16384 __ __