forked from phoenix-oss/llama-stack-mirror
[inference] Add a TGI adapter (#52)
* TGI adapter and some refactoring of other inference adapters * Use the lower-level `generate_stream()` method for correct tool calling --------- Co-authored-by: Ashwin Bharambe <ashwin@meta.com>
This commit is contained in:
parent
6ad7365676
commit
21bedc1596
3 changed files with 256 additions and 0 deletions
15
llama_toolchain/inference/adapters/tgi/__init__.py
Normal file
15
llama_toolchain/inference/adapters/tgi/__init__.py
Normal file
|
@ -0,0 +1,15 @@
|
|||
# Copyright (c) Meta Platforms, Inc. and affiliates.
|
||||
# All rights reserved.
|
||||
#
|
||||
# This source code is licensed under the terms described in the LICENSE file in
|
||||
# the root directory of this source tree.
|
||||
|
||||
from llama_toolchain.core.datatypes import RemoteProviderConfig
|
||||
|
||||
|
||||
async def get_adapter_impl(config: RemoteProviderConfig, _deps):
|
||||
from .tgi import TGIInferenceAdapter
|
||||
|
||||
impl = TGIInferenceAdapter(config.url)
|
||||
await impl.initialize()
|
||||
return impl
|
Loading…
Add table
Add a link
Reference in a new issue