Attach an adapter
Attach a LoRA adapter model to a serving_mode=‘multi_adapter’ deployment under a stable public alias (serve_name). The fleet’s replicas download and load the adapter asynchronously: the attachment starts in status ‘attaching’ and becomes ‘active’ once every serving replica reports it loaded (poll the list endpoint).
curl --request POST \
--url https://api.veri.studio/v1/deployments/{deployment_id}/adapters \
--header 'Authorization: Bearer <token>' \
--header 'Content-Type: application/json' \
--data '
{
"serve_name": "<string>",
"model_id": "<string>"
}
'import requests
url = "https://api.veri.studio/v1/deployments/{deployment_id}/adapters"
payload = {
"serve_name": "<string>",
"model_id": "<string>"
}
headers = {
"Authorization": "Bearer <token>",
"Content-Type": "application/json"
}
response = requests.post(url, json=payload, headers=headers)
print(response.text)const options = {
method: 'POST',
headers: {Authorization: 'Bearer <token>', 'Content-Type': 'application/json'},
body: JSON.stringify({serve_name: '<string>', model_id: '<string>'})
};
fetch('https://api.veri.studio/v1/deployments/{deployment_id}/adapters', options)
.then(res => res.json())
.then(res => console.log(res))
.catch(err => console.error(err));{
"object": "<string>",
"id": "<string>",
"deployment_id": "<string>",
"serve_name": "<string>",
"model_id": "<string>",
"version": 123,
"internal_name": "<string>",
"status": "<string>",
"created_at": "2023-11-07T05:31:56Z",
"updated_at": "2023-11-07T05:31:56Z",
"replicas": [
{
"replica_id": "<string>",
"state": "<string>",
"error": "<string>"
}
],
"previous_model_id": "<string>",
"pending_model_id": "<string>",
"pending_internal_name": "<string>",
"model_version_id": "<string>",
"pending_model_version_id": "<string>"
}Authorizations
API key with the vk_ prefix. Create one from the dashboard.
Path Parameters
Deployment ID
Body
Public routing alias for the adapter: 1-64 chars of [A-Za-z0-9._-]. Requests to the deployment addressed with this model name are served by the adapter. Must not collide with the deployment's base alias.
ID of a ready custom model with artifact_type='adapter' whose base_model matches the deployment's base model.
Response
The new attachment, status 'attaching'; replicas[] carries the per-replica load state
A LoRA adapter attached to a multi_adapter deployment.
Always "deployment_adapter".
Attachment ID (ada_...).
The deployment (fleet) this adapter is attached to.
Public routing alias the adapter serves under.
The custom model currently served under this alias.
Monotonic version counter; incremented on every flip/rollback.
vLLM-facing name ("{serve_name}@v{version}-{tag}", tag derived from model_id): never reused for different weights.
attaching | active | detaching. 'active' means every serving replica reports the current version loaded.
Attachment creation time.
Last status/version change time.
Per-serving-replica load state for the CURRENT internal_name.
Show child attributes
Show child attributes
The previously served model, if any (the rollback target).
Mid-flip only: the staged new version awaiting fleet convergence.
Mid-flip only: the staged version's internal name.
The platform model version (model_versions.id) currently served under this alias, when the model belongs to a lineage.
Mid-flip only: the platform model version staged by a production
alias move; the fleet cuts over to it automatically on convergence.
curl --request POST \
--url https://api.veri.studio/v1/deployments/{deployment_id}/adapters \
--header 'Authorization: Bearer <token>' \
--header 'Content-Type: application/json' \
--data '
{
"serve_name": "<string>",
"model_id": "<string>"
}
'import requests
url = "https://api.veri.studio/v1/deployments/{deployment_id}/adapters"
payload = {
"serve_name": "<string>",
"model_id": "<string>"
}
headers = {
"Authorization": "Bearer <token>",
"Content-Type": "application/json"
}
response = requests.post(url, json=payload, headers=headers)
print(response.text)const options = {
method: 'POST',
headers: {Authorization: 'Bearer <token>', 'Content-Type': 'application/json'},
body: JSON.stringify({serve_name: '<string>', model_id: '<string>'})
};
fetch('https://api.veri.studio/v1/deployments/{deployment_id}/adapters', options)
.then(res => res.json())
.then(res => console.log(res))
.catch(err => console.error(err));{
"object": "<string>",
"id": "<string>",
"deployment_id": "<string>",
"serve_name": "<string>",
"model_id": "<string>",
"version": 123,
"internal_name": "<string>",
"status": "<string>",
"created_at": "2023-11-07T05:31:56Z",
"updated_at": "2023-11-07T05:31:56Z",
"replicas": [
{
"replica_id": "<string>",
"state": "<string>",
"error": "<string>"
}
],
"previous_model_id": "<string>",
"pending_model_id": "<string>",
"pending_internal_name": "<string>",
"model_version_id": "<string>",
"pending_model_version_id": "<string>"
}
