Roboflow Model API

The Model API from Roboflow — 1 operation(s) for model.

OpenAPI Specification

roboflow-model-api-openapi.yml Raw ↑
openapi: 3.2.0
info:
  title: Roboflow Inference Server Model API
  description: Roboflow inference server
  termsOfService: https://roboflow.com/terms
  contact:
    name: Roboflow Inc.
    url: https://roboflow.com/contact
    email: help@roboflow.com
  license:
    name: Apache 2.0
    url: https://www.apache.org/licenses/LICENSE-2.0.html
  version: 1.3.8
tags:
- name: Model
paths:
  /model/registry:
    get:
      summary: Get model keys
      description: Get the ID of each loaded model
      operationId: registry_model_registry_get
      responses:
        '200':
          description: Successful Response
          content:
            application/json:
              schema:
                $ref: '#/components/schemas/ModelsDescriptions'
      tags:
      - Model
components:
  schemas:
    ModelsDescriptions:
      properties:
        models:
          items:
            $ref: '#/components/schemas/ModelDescriptionEntity'
          type: array
          title: Models
          description: List of models that are loaded by model manager.
        total_vram_bytes:
          anyOf:
          - type: integer
          - type: 'null'
          title: Total Vram Bytes
          description: Total estimated VRAM consumed by all loaded models in bytes.
        gpu_memory_used:
          anyOf:
          - type: integer
          - type: 'null'
          title: Gpu Memory Used
          description: Current GPU memory in use in bytes (device-level, includes all runtimes).
        gpu_memory_total:
          anyOf:
          - type: integer
          - type: 'null'
          title: Gpu Memory Total
          description: Total GPU memory available in bytes.
        torch_cuda_allocated:
          anyOf:
          - type: integer
          - type: 'null'
          title: Torch Cuda Allocated
          description: Live tensor memory allocated by PyTorch's CUDA allocator in bytes.
        torch_cuda_reserved:
          anyOf:
          - type: integer
          - type: 'null'
          title: Torch Cuda Reserved
          description: Total memory reserved by PyTorch's CUDA allocator in bytes.
        torch_cuda_allocator_cache:
          anyOf:
          - type: integer
          - type: 'null'
          title: Torch Cuda Allocator Cache
          description: Reserved but currently unallocated PyTorch CUDA memory in bytes.
        non_torch_gpu_memory:
          anyOf:
          - type: integer
          - type: 'null'
          title: Non Torch Gpu Memory
          description: Device memory not reserved by PyTorch in bytes. This includes native runtimes,
            CUDA context overhead, and allocations from other processes.
      type: object
      required:
      - models
      title: ModelsDescriptions
    ModelDescriptionEntity:
      properties:
        model_id:
          type: string
          title: Model Id
          description: Identifier of the model
          examples:
          - some-project/3
        task_type:
          type: string
          title: Task Type
          description: Type of the task that the model performs
          examples:
          - classification
        batch_size:
          anyOf:
          - type: integer
          - type: 'null'
          title: Batch Size
          description: Batch size accepted by the model (if registered).
        input_height:
          anyOf:
          - type: integer
          - type: 'null'
          title: Input Height
          description: Image input height accepted by the model (if registered).
        input_width:
          anyOf:
          - type: integer
          - type: 'null'
          title: Input Width
          description: Image input width accepted by the model (if registered).
        vram_bytes:
          anyOf:
          - type: integer
          - type: 'null'
          title: Vram Bytes
          description: Estimated GPU VRAM consumed by this model in bytes (measured during load).
        request_aliases:
          items:
            type: string
          type: array
          title: Request Aliases
          description: Other model IDs that resolved to this model.
        request_paths:
          items:
            type: string
          type: array
          title: Request Paths
          description: HTTP request paths that triggered inference on this model (e.g. /door-glyph-locator/10,
            /infer/object_detection).
      type: object
      required:
      - model_id
      - task_type
      title: ModelDescriptionEntity