pub struct DedicatedInferencesPatchResponseDedicatedInference {
pub created_at: Option<DateTime<Utc>>,
pub endpoints: Option<DedicatedInferencesPatchResponseDedicatedInferenceEndpoints>,
pub id: Option<Uuid>,
pub pending_deployment_spec: Option<DedicatedInferencesPatchResponseDedicatedInferencePendingDeploymentSpec>,
pub region: Option<String>,
pub spec: Option<DedicatedInferencesPatchResponseDedicatedInferenceSpec>,
pub status: Option<DedicatedInferencesPatchResponseDedicatedInferenceStatus>,
pub updated_at: Option<DateTime<Utc>>,
pub vpc_uuid: Option<Uuid>,
}Expand description
A Dedicated Inference instance.
JSON schema
{
"description": "A Dedicated Inference instance.",
"type": "object",
"properties": {
"created_at": {
"description": "When the Dedicated Inference was created.",
"readOnly": true,
"examples": [
"2024-01-09T20:44:32Z"
],
"type": "string",
"format": "date-time"
},
"endpoints": {
"readOnly": true,
"type": "object",
"properties": {
"private_endpoint_fqdn": {
"description": "Private VPC FQDN of the Dedicated Inference instance.",
"examples": [
"https://b4bfug4jc41kts2ro54if91eo-private-dedicated-inference.do-infra.ai"
],
"type": "string",
"format": "uri"
},
"public_endpoint_fqdn": {
"description": "Public FQDN of the Dedicated Inference instance.",
"examples": [
"https://b4bfug4jc41kts2ro54if91eo-public-dedicated-inference.do-infra.ai"
],
"type": "string",
"format": "uri"
}
}
},
"id": {
"description": "Unique ID of the Dedicated Inference.",
"readOnly": true,
"examples": [
"6b5c619c-359c-44ca-87e2-47e98170c01d"
],
"type": "string",
"format": "uuid"
},
"pending_deployment_spec": {
"description": "Pending deployment when status is provisioning or updating.",
"type": "object",
"properties": {
"created_at": {
"examples": [
"2024-01-09T20:44:32Z"
],
"type": "string",
"format": "date-time"
},
"enable_public_endpoint": {
"description": "Whether to expose a public LLM endpoint.",
"type": "boolean"
},
"id": {
"description": "Deployment UUID.",
"examples": [
"7c6d729d-360d-44db-88e3-58e98281d12e"
],
"type": "string",
"format": "uuid"
},
"model_deployments": {
"description": "At least one model deployment is required.",
"type": "array",
"items": {
"description": "Fallback schema for missing file reference",
"type": "object",
"additionalProperties": true
},
"minItems": 1
},
"name": {
"description": "Name of the Dedicated Inference. Must be unique within the team.",
"examples": [
"new-dedicated-inference"
],
"type": "string",
"maxLength": 255,
"pattern": "^[a-zA-Z0-9_-]+$"
},
"status": {
"examples": [
"provisioning"
],
"type": "string",
"enum": [
"provisioning",
"updating"
]
},
"updated_at": {
"examples": [
"2024-01-09T20:44:32Z"
],
"type": "string",
"format": "date-time"
},
"version": {
"description": "Spec version.",
"examples": [
1
],
"type": "integer"
},
"vpc": {
"type": "object",
"required": [
"uuid"
],
"properties": {
"uuid": {
"description": "VPC UUID for the Dedicated Inference.",
"examples": [
"997615ce-132d-4bae-9270-9ee21b395e5d"
],
"type": "string",
"format": "uuid"
}
}
}
}
},
"region": {
"description": "DigitalOcean region where the Dedicated Inference is hosted.",
"readOnly": true,
"examples": [
"atl1"
],
"type": "string"
},
"spec": {
"description": "Structured configuration for a Dedicated Inference deployment.",
"type": "object",
"required": [
"enable_public_endpoint",
"model_deployments",
"name",
"region",
"version",
"vpc"
],
"properties": {
"enable_public_endpoint": {
"description": "Whether to expose a public LLM endpoint.",
"type": "boolean"
},
"model_deployments": {
"description": "At least one model deployment is required.",
"type": "array",
"items": {
"description": "Configuration for a single model deployment.",
"type": "object",
"properties": {
"accelerators": {
"description": "Accelerator configuration for this deployment.",
"type": "array",
"items": {
"description": "Auto-generated fallback definition for: accelerator_config_spec",
"type": "object",
"additionalProperties": true
}
},
"model_id": {
"description": "Used to identify an existing deployment when updating; empty means create new.",
"examples": [
""
],
"type": "string"
},
"model_provider": {
"description": "Model provider.",
"examples": [
"hugging_face"
],
"type": "string",
"enum": [
"hugging_face"
]
},
"model_slug": {
"description": "Model identifier (e.g. Hugging Face slug).",
"examples": [
"mistral/mistral-7b-instruct-v3"
],
"type": "string"
},
"workload_config": {
"description": "Workload-specific configuration (e.g. ISL/OSL in future).",
"type": "object"
}
}
},
"minItems": 1
},
"name": {
"description": "Name of the Dedicated Inference. Must be unique within the team.",
"examples": [
"new-dedicated-inference"
],
"type": "string",
"maxLength": 255,
"pattern": "^[a-zA-Z0-9_-]+$"
},
"region": {
"description": "DigitalOcean region where the Dedicated Inference is hosted.",
"examples": [
"atl1"
],
"type": "string",
"enum": [
"atl1",
"nyc2",
"tor1"
]
},
"version": {
"description": "Spec version.",
"examples": [
1
],
"type": "integer"
},
"vpc": {
"type": "object",
"required": [
"uuid"
],
"properties": {
"uuid": {
"description": "VPC UUID for the Dedicated Inference.",
"examples": [
"997615ce-132d-4bae-9270-9ee21b395e5d"
],
"type": "string",
"format": "uuid"
}
}
}
}
},
"status": {
"description": "Current state of the Dedicated Inference.",
"readOnly": true,
"examples": [
"active"
],
"type": "string",
"enum": [
"active",
"new",
"provisioning",
"updating",
"deleting",
"error"
]
},
"updated_at": {
"description": "When the Dedicated Inference was last updated.",
"readOnly": true,
"examples": [
"2024-01-09T20:44:32Z"
],
"type": "string",
"format": "date-time"
},
"vpc_uuid": {
"description": "VPC UUID of the Dedicated Inference.",
"readOnly": true,
"examples": [
"997615ce-132d-4bae-9270-9ee21b395e5d"
],
"type": "string",
"format": "uuid"
}
}
}Fields§
§created_at: Option<DateTime<Utc>>When the Dedicated Inference was created.
endpoints: Option<DedicatedInferencesPatchResponseDedicatedInferenceEndpoints>§id: Option<Uuid>Unique ID of the Dedicated Inference.
pending_deployment_spec: Option<DedicatedInferencesPatchResponseDedicatedInferencePendingDeploymentSpec>§region: Option<String>DigitalOcean region where the Dedicated Inference is hosted.
spec: Option<DedicatedInferencesPatchResponseDedicatedInferenceSpec>§status: Option<DedicatedInferencesPatchResponseDedicatedInferenceStatus>Current state of the Dedicated Inference.
updated_at: Option<DateTime<Utc>>When the Dedicated Inference was last updated.
vpc_uuid: Option<Uuid>VPC UUID of the Dedicated Inference.
Trait Implementations§
Source§impl Clone for DedicatedInferencesPatchResponseDedicatedInference
impl Clone for DedicatedInferencesPatchResponseDedicatedInference
Source§fn clone(&self) -> DedicatedInferencesPatchResponseDedicatedInference
fn clone(&self) -> DedicatedInferencesPatchResponseDedicatedInference
Returns a duplicate of the value. Read more
1.0.0 (const: unstable) · Source§fn clone_from(&mut self, source: &Self)
fn clone_from(&mut self, source: &Self)
Performs copy-assignment from
source. Read moreSource§impl<'de> Deserialize<'de> for DedicatedInferencesPatchResponseDedicatedInference
impl<'de> Deserialize<'de> for DedicatedInferencesPatchResponseDedicatedInference
Source§fn deserialize<__D>(__deserializer: __D) -> Result<Self, __D::Error>where
__D: Deserializer<'de>,
fn deserialize<__D>(__deserializer: __D) -> Result<Self, __D::Error>where
__D: Deserializer<'de>,
Deserialize this value from the given Serde deserializer. Read more
Source§impl From<&DedicatedInferencesPatchResponseDedicatedInference> for DedicatedInferencesPatchResponseDedicatedInference
impl From<&DedicatedInferencesPatchResponseDedicatedInference> for DedicatedInferencesPatchResponseDedicatedInference
Source§fn from(value: &DedicatedInferencesPatchResponseDedicatedInference) -> Self
fn from(value: &DedicatedInferencesPatchResponseDedicatedInference) -> Self
Converts to this type from the input type.
Auto Trait Implementations§
impl Freeze for DedicatedInferencesPatchResponseDedicatedInference
impl RefUnwindSafe for DedicatedInferencesPatchResponseDedicatedInference
impl Send for DedicatedInferencesPatchResponseDedicatedInference
impl Sync for DedicatedInferencesPatchResponseDedicatedInference
impl Unpin for DedicatedInferencesPatchResponseDedicatedInference
impl UnsafeUnpin for DedicatedInferencesPatchResponseDedicatedInference
impl UnwindSafe for DedicatedInferencesPatchResponseDedicatedInference
Blanket Implementations§
Source§impl<T> BorrowMut<T> for Twhere
T: ?Sized,
impl<T> BorrowMut<T> for Twhere
T: ?Sized,
Source§fn borrow_mut(&mut self) -> &mut T
fn borrow_mut(&mut self) -> &mut T
Mutably borrows from an owned value. Read more