docker run --rm --gpus all -p 8000:8000 \ vllm/vllm-openai:latest \ --model mistralai/Mistral-7B-v0.1 \ --gpu-memory-utilization 0.85 __ __