2023-08-10 08:24:52 +00:00
|
|
|
- sections:
|
|
|
|
- local: index
|
|
|
|
title: Text Generation Inference
|
|
|
|
- local: quicktour
|
|
|
|
title: Quick Tour
|
|
|
|
- local: installation
|
|
|
|
title: Installation
|
|
|
|
- local: supported_models
|
|
|
|
title: Supported Models and Hardware
|
2024-01-24 16:41:28 +00:00
|
|
|
- local: messages_api
|
|
|
|
title: Messages API
|
2023-08-10 08:24:52 +00:00
|
|
|
title: Getting started
|
|
|
|
- sections:
|
|
|
|
- local: basic_tutorials/consuming_tgi
|
|
|
|
title: Consuming TGI
|
|
|
|
- local: basic_tutorials/preparing_model
|
|
|
|
title: Preparing Model for Serving
|
|
|
|
- local: basic_tutorials/gated_model_access
|
|
|
|
title: Serving Private & Gated Models
|
2023-08-10 13:00:30 +00:00
|
|
|
- local: basic_tutorials/using_cli
|
|
|
|
title: Using TGI CLI
|
2023-09-27 14:01:38 +00:00
|
|
|
- local: basic_tutorials/launcher
|
|
|
|
title: All TGI CLI options
|
2023-09-12 13:55:14 +00:00
|
|
|
- local: basic_tutorials/non_core_models
|
|
|
|
title: Non-core Model Serving
|
2024-04-05 11:32:53 +00:00
|
|
|
- local: basic_tutorials/safety
|
|
|
|
title: Safety
|
2024-04-30 10:14:39 +00:00
|
|
|
- local: basic_tutorials/visual_language_models
|
|
|
|
title: Visual Language Models
|
2023-08-10 08:24:52 +00:00
|
|
|
title: Tutorials
|
2023-08-18 11:27:08 +00:00
|
|
|
- sections:
|
|
|
|
- local: conceptual/streaming
|
|
|
|
title: Streaming
|
2023-09-12 13:52:46 +00:00
|
|
|
- local: conceptual/quantization
|
|
|
|
title: Quantization
|
2023-09-12 10:11:20 +00:00
|
|
|
- local: conceptual/tensor_parallelism
|
|
|
|
title: Tensor Parallelism
|
2023-09-08 12:18:42 +00:00
|
|
|
- local: conceptual/paged_attention
|
|
|
|
title: PagedAttention
|
2023-09-07 14:22:06 +00:00
|
|
|
- local: conceptual/safetensors
|
|
|
|
title: Safetensors
|
2023-09-06 13:36:49 +00:00
|
|
|
- local: conceptual/flash_attention
|
|
|
|
title: Flash Attention
|
2024-02-28 10:30:37 +00:00
|
|
|
- local: conceptual/speculation
|
|
|
|
title: Speculation (Medusa, ngram)
|
|
|
|
- local: conceptual/guidance
|
|
|
|
title: Guidance, JSON, tools (using outlines)
|
2024-04-30 10:14:39 +00:00
|
|
|
|
2023-08-18 11:27:08 +00:00
|
|
|
title: Conceptual Guides
|