Roll back an adapter
Point the alias back at the previously served adapter model (previous_model_id) under a fresh version. Status returns to ‘attaching’ until every serving replica reports the rollback target loaded: usually immediate when the old weights are still resident, otherwise after a re-download.
curl --request POST \
--url https://api.veri.studio/v1/deployments/{deployment_id}/adapters/{serve_name}/rollback \
--header 'Authorization: Bearer <token>'import requests
url = "https://api.veri.studio/v1/deployments/{deployment_id}/adapters/{serve_name}/rollback"
headers = {"Authorization": "Bearer <token>"}
response = requests.post(url, headers=headers)
print(response.text)const options = {method: 'POST', headers: {Authorization: 'Bearer <token>'}};
fetch('https://api.veri.studio/v1/deployments/{deployment_id}/adapters/{serve_name}/rollback', options)
.then(res => res.json())
.then(res => console.log(res))
.catch(err => console.error(err));{
"object": "<string>",
"id": "<string>",
"deployment_id": "<string>",
"serve_name": "<string>",
"model_id": "<string>",
"version": 123,
"internal_name": "<string>",
"status": "<string>",
"created_at": "2023-11-07T05:31:56Z",
"updated_at": "2023-11-07T05:31:56Z",
"replicas": [
{
"replica_id": "<string>",
"state": "<string>",
"error": "<string>"
}
],
"previous_model_id": "<string>",
"pending_model_id": "<string>",
"pending_internal_name": "<string>",
"model_version_id": "<string>",
"pending_model_version_id": "<string>"
}Authorizations
API key with the vk_ prefix. Create one from the dashboard.
Path Parameters
Deployment ID
The adapter's public alias
Response
The alias now serves the previous model at version+1, status 'attaching' until the fleet converges
A LoRA adapter attached to a multi_adapter deployment.
Always "deployment_adapter".
Attachment ID (ada_...).
The deployment (fleet) this adapter is attached to.
Public routing alias the adapter serves under.
The custom model currently served under this alias.
Monotonic version counter; incremented on every flip/rollback.
vLLM-facing name ("{serve_name}@v{version}-{tag}", tag derived from model_id): never reused for different weights.
attaching | active | detaching. 'active' means every serving replica reports the current version loaded.
Attachment creation time.
Last status/version change time.
Per-serving-replica load state for the CURRENT internal_name.
Show child attributes
Show child attributes
The previously served model, if any (the rollback target).
Mid-flip only: the staged new version awaiting fleet convergence.
Mid-flip only: the staged version's internal name.
The platform model version (model_versions.id) currently served under this alias, when the model belongs to a lineage.
Mid-flip only: the platform model version staged by a production
alias move; the fleet cuts over to it automatically on convergence.
curl --request POST \
--url https://api.veri.studio/v1/deployments/{deployment_id}/adapters/{serve_name}/rollback \
--header 'Authorization: Bearer <token>'import requests
url = "https://api.veri.studio/v1/deployments/{deployment_id}/adapters/{serve_name}/rollback"
headers = {"Authorization": "Bearer <token>"}
response = requests.post(url, headers=headers)
print(response.text)const options = {method: 'POST', headers: {Authorization: 'Bearer <token>'}};
fetch('https://api.veri.studio/v1/deployments/{deployment_id}/adapters/{serve_name}/rollback', options)
.then(res => res.json())
.then(res => console.log(res))
.catch(err => console.error(err));{
"object": "<string>",
"id": "<string>",
"deployment_id": "<string>",
"serve_name": "<string>",
"model_id": "<string>",
"version": 123,
"internal_name": "<string>",
"status": "<string>",
"created_at": "2023-11-07T05:31:56Z",
"updated_at": "2023-11-07T05:31:56Z",
"replicas": [
{
"replica_id": "<string>",
"state": "<string>",
"error": "<string>"
}
],
"previous_model_id": "<string>",
"pending_model_id": "<string>",
"pending_internal_name": "<string>",
"model_version_id": "<string>",
"pending_model_version_id": "<string>"
}
