Skip to main content

TokenTrackingLLM

Struct TokenTrackingLLM 

Source
pub struct TokenTrackingLLM<L: BaseChatModel> { /* private fields */ }
Expand description

LLM wrapper with token statistics

Wraps any BaseChatModel, accumulating prompt / completion token usage automatically, preferring the real usage returned by the LLM, falling back to tiktoken estimates.

v0.20.1: implements BaseChatModel itself, so a tracked LLM can be plugged directly into framework agents (e.g. FunctionCallingAgent::new); every chat/stream/bind_tools call inside the agent loop is counted, and get_usage/estimate_cost reflect the cumulative agent-run usage.

Implementations§

Source§

impl<L: BaseChatModel> TokenTrackingLLM<L>

Source

pub fn new(llm: L, counter: Arc<dyn TokenCounter>) -> Self

Wraps an LLM with a custom counter.

Source

pub fn for_openai(llm: L) -> Result<Self, TokenCounterError>

Wraps with a Tiktoken (cl100k_base) counter

Source

pub fn with_metrics_sink(self, sink: Arc<dyn MetricsSink>) -> Self

Attaches an observability sink: after each call that reports usage the wrapper exports a TokenUsage event (real when the provider reports it, otherwise the tiktoken estimate). The sink is shared across wrappers rebuilt by bind_tools/with_temperature/with_max_tokens.

Source

pub async fn chat( &self, messages: Vec<Message>, config: Option<RunnableConfig>, ) -> Result<LLMResult, L::Error>

Calls the LLM and counts tokens.

Inherent method (kept for backward compatibility); delegates to Self::chat_tracked. When the model is reached through dyn BaseChatModel (e.g. inside an agent), the trait chat takes over — both count through the same helper, so the two paths never diverge.

Source

pub async fn get_usage(&self) -> TrackerTokenUsage

Returns the cumulative usage

Source

pub async fn reset(&self)

Resets the statistics

Source

pub async fn estimate_cost(&self, pricing: &ModelPricing) -> f64

Estimates the cost (USD)

Trait Implementations§

Source§

impl<L> BaseChatModel for TokenTrackingLLM<L>
where L: BaseChatModel + Send + Sync,

Source§

fn chat<'life0, 'async_trait>( &'life0 self, messages: Vec<Message>, config: Option<RunnableConfig>, ) -> Pin<Box<dyn Future<Output = Result<LLMResult, Self::Error>> + Send + 'async_trait>>
where Self: 'async_trait, 'life0: 'async_trait,

Chat with the model. Read more
Source§

fn stream_chat<'life0, 'async_trait>( &'life0 self, messages: Vec<Message>, config: Option<RunnableConfig>, ) -> Pin<Box<dyn Future<Output = Result<Pin<Box<dyn Stream<Item = Result<StreamChunk, Self::Error>> + Send>>, Self::Error>> + Send + 'async_trait>>
where Self: 'async_trait, 'life0: 'async_trait,

Stream chat with the model. Read more
Source§

fn bind_tools( &self, tools: Vec<ToolDefinition>, ) -> Option<Box<dyn BaseChatModel<Error = Self::Error> + Send + Sync>>

Bind tool definitions for function calling. Read more
Source§

fn chat_with_system<'life0, 'async_trait>( &'life0 self, system: String, messages: Vec<Message>, ) -> Pin<Box<dyn Future<Output = Result<LLMResult, Self::Error>> + Send + 'async_trait>>
where Self: Sync + 'async_trait, 'life0: 'async_trait,

Chat with system prompt. Read more
Source§

impl<L> BaseLanguageModel<Vec<Message>, LLMResult> for TokenTrackingLLM<L>
where L: BaseChatModel + Send + Sync,

Source§

fn model_name(&self) -> &str

Returns the model name.
Source§

fn get_num_tokens(&self, text: &str) -> usize

Calculates token count for text. Read more
Source§

fn temperature(&self) -> Option<f32>

Returns the temperature parameter.
Source§

fn max_tokens(&self) -> Option<usize>

Returns the max tokens limit.
Source§

fn with_temperature(self, temp: f32) -> Self
where Self: Sized,

Sets the temperature parameter.
Source§

fn with_max_tokens(self, max: usize) -> Self
where Self: Sized,

Sets the max tokens limit.
Source§

impl<L> Runnable<Vec<Message>, LLMResult> for TokenTrackingLLM<L>
where L: BaseChatModel + Send + Sync,

Source§

type Error = <L as Runnable<Vec<Message>, LLMResult>>::Error

Error type.
Source§

fn invoke<'life0, 'async_trait>( &'life0 self, input: Vec<Message>, config: Option<RunnableConfig>, ) -> Pin<Box<dyn Future<Output = Result<LLMResult, Self::Error>> + Send + 'async_trait>>
where Self: 'async_trait, 'life0: 'async_trait,

Transforms single input to output. Read more
Source§

fn stream<'life0, 'async_trait>( &'life0 self, input: Vec<Message>, config: Option<RunnableConfig>, ) -> Pin<Box<dyn Future<Output = Result<Pin<Box<dyn Stream<Item = Result<LLMResult, Self::Error>> + Send>>, Self::Error>> + Send + 'async_trait>>
where Self: 'async_trait, 'life0: 'async_trait,

Streaming output - for real-time responses (LLM, etc). Read more
Source§

fn batch<'life0, 'async_trait>( &'life0 self, inputs: Vec<Input>, config: Option<RunnableConfig>, ) -> Pin<Box<dyn Future<Output = Result<Vec<Output>, Self::Error>> + Send + 'async_trait>>
where Self: 'async_trait, 'life0: 'async_trait,

Batch processing - transforms multiple inputs to outputs. Read more
Source§

fn batch_as_completed<'life0, 'async_trait>( &'life0 self, inputs: Vec<Input>, config: Option<RunnableConfig>, ) -> Pin<Box<dyn Future<Output = Result<Vec<(usize, Output)>, Self::Error>> + Send + 'async_trait>>
where Self: 'async_trait, 'life0: 'async_trait,

Batch processing that returns results in completion order. Read more
Source§

fn transform<'life0, 'async_trait>( &'life0 self, input: Pin<Box<dyn Stream<Item = Result<Input, Self::Error>> + Send>>, config: Option<RunnableConfig>, ) -> Pin<Box<dyn Future<Output = Result<Pin<Box<dyn Stream<Item = Result<Output, Self::Error>> + Send + '_>>, Self::Error>> + Send + 'async_trait>>
where Self: 'async_trait, 'life0: 'async_trait,

Stream-to-stream transformation - the core of LCEL streaming. Read more

Auto Trait Implementations§

§

impl<L> !RefUnwindSafe for TokenTrackingLLM<L>

§

impl<L> !UnwindSafe for TokenTrackingLLM<L>

§

impl<L> Freeze for TokenTrackingLLM<L>
where L: Freeze,

§

impl<L> Send for TokenTrackingLLM<L>

§

impl<L> Sync for TokenTrackingLLM<L>

§

impl<L> Unpin for TokenTrackingLLM<L>
where L: Unpin,

§

impl<L> UnsafeUnpin for TokenTrackingLLM<L>
where L: UnsafeUnpin,

Blanket Implementations§

Source§

impl<T> Any for T
where T: 'static + ?Sized,

Source§

fn type_id(&self) -> TypeId

Gets the TypeId of self. Read more
Source§

impl<T> Borrow<T> for T
where T: ?Sized,

Source§

fn borrow(&self) -> &T

Immutably borrows from an owned value. Read more
Source§

impl<T> BorrowMut<T> for T
where T: ?Sized,

Source§

fn borrow_mut(&mut self) -> &mut T

Mutably borrows from an owned value. Read more
Source§

impl<T> From<T> for T

Source§

fn from(t: T) -> T

Returns the argument unchanged.

Source§

impl<T> Instrument for T

Source§

fn instrument(self, span: Span) -> Instrumented<Self>

Instruments this type with the provided Span, returning an Instrumented wrapper. Read more
Source§

fn in_current_span(self) -> Instrumented<Self>

Instruments this type with the current Span, returning an Instrumented wrapper. Read more
Source§

impl<T, U> Into<U> for T
where U: From<T>,

Source§

fn into(self) -> U

Calls U::from(self).

That is, this conversion is whatever the implementation of From<T> for U chooses to do.

Source§

impl<T> PolicyExt for T
where T: ?Sized,

Source§

fn and<P, B, E>(self, other: P) -> And<T, P>
where T: Sized + Policy<B, E>, P: Policy<B, E>,

Create a new Policy that returns Action::Follow only if self and other return Action::Follow. Read more
Source§

fn or<P, B, E>(self, other: P) -> Or<T, P>
where T: Sized + Policy<B, E>, P: Policy<B, E>,

Create a new Policy that returns Action::Follow if either self or other returns Action::Follow. Read more
Source§

impl<I, O, R> RunnableExt<I, O> for R
where I: Send + Sync + 'static, O: Send + Sync + 'static, R: Runnable<I, O> + 'static, <R as Runnable<I, O>>::Error: Into<LcelError>,

Source§

fn pipe<O2, R2>(self, other: R2) -> RunnableSequence<Input, O2>
where O2: Send + Sync + 'static, R2: Runnable<Output, O2> + Send + Sync + 'static, R2::Error: Into<LcelError>,

Pipe the output of this runnable into another runnable. Read more
Source§

fn into_sequence(self) -> RunnableSequence<Input, Output>

Create a RunnableSequence from this runnable as a single step. Read more
Source§

fn with_fallbacks<R>( self, fallbacks: Vec<R>, ) -> RunnableWithFallbacks<Input, Output>
where Input: Clone, R: Runnable<Input, Output> + Send + Sync + 'static, R::Error: Into<LcelError>,

Add fallback runnables that are tried if this one fails. Read more
Source§

fn with_retry(self, retry_config: RetryConfig) -> RunnableRetry<Input, Output>
where Input: Clone,

Wrap this runnable with retry logic using exponential backoff. Read more
Source§

fn configurable_alternatives<K, R>( self, which: impl Into<String>, default_key: impl Into<String>, alternatives: Vec<(K, R)>, ) -> RunnableConfigurable<Input, Output>
where K: Into<String>, R: Runnable<Input, Output> + Send + Sync + 'static, R::Error: Into<LcelError>,

Route between a default runnable and named alternatives at invoke time. Read more
Source§

fn configurable_fields(self) -> RunnableConfigurableFields<Input, Output>

Override recognized config fields at invoke time from config.configurable (Python’s Runnable.configurable_fields). Read more
Source§

impl<M> StreamingStructuredOutputExt for M
where M: BaseChatModel,

Source§

fn stream_structured_output<'life0, 'life1, 'async_trait, T>( &'life0 self, schema: Value, prompt: &'life1 str, ) -> Pin<Box<dyn Future<Output = Result<Pin<Box<dyn Stream<Item = Result<T, StructuredOutputError>> + Send>>, StructuredOutputError>> + Send + 'async_trait>>
where T: DeserializeOwned + Serialize + Clone + PartialEq + Unpin + Send + Sync + 'static + 'async_trait, Self: Sync + 'async_trait, 'life0: 'async_trait, 'life1: 'async_trait,

Stream structured output from a chat model. Read more
Source§

impl<M> StructuredOutputExt for M
where M: BaseChatModel,

Source§

fn with_structured_output<'life0, 'life1, 'async_trait, T>( &'life0 self, schema: Value, prompt: &'life1 str, ) -> Pin<Box<dyn Future<Output = Result<T, StructuredOutputError>> + Send + 'async_trait>>
where T: 'async_trait + DeserializeOwned + Serialize + Send + Sync + 'static, Self: Sync + 'async_trait, 'life0: 'async_trait, 'life1: 'async_trait,

Call the LLM with a JSON schema and prompt, returning a parsed result of type T. Read more
Source§

impl<T, U> TryFrom<U> for T
where U: Into<T>,

Source§

type Error = !

The type returned in the event of a conversion error.
Source§

fn try_from(value: U) -> Result<T, <T as TryFrom<U>>::Error>

Performs the conversion.
Source§

impl<T, U> TryInto<U> for T
where U: TryFrom<T>,

Source§

type Error = <U as TryFrom<T>>::Error

The type returned in the event of a conversion error.
Source§

fn try_into(self) -> Result<U, <U as TryFrom<T>>::Error>

Performs the conversion.
Source§

impl<T> WithSubscriber for T

Source§

fn with_subscriber<S>(self, subscriber: S) -> WithDispatch<Self>
where S: Into<Dispatch>,

Attaches the provided Subscriber to this type, returning a WithDispatch wrapper. Read more
Source§

fn with_current_subscriber(self) -> WithDispatch<Self>

Attaches the current default Subscriber to this type, returning a WithDispatch wrapper. Read more