Add support for prompt variables (bot name, date/time, model id)
Fixes https://github.com/etkecc/baibot/issues/10 This also includes them in the default prompts (for newly-created agents), so that people can get a better experience out of the box.
This commit is contained in:
10
Cargo.lock
generated
10
Cargo.lock
generated
@@ -303,6 +303,7 @@ dependencies = [
|
|||||||
"anyhow",
|
"anyhow",
|
||||||
"async-openai 0.24.0",
|
"async-openai 0.24.0",
|
||||||
"base64 0.22.1",
|
"base64 0.22.1",
|
||||||
|
"chrono",
|
||||||
"etke_openai_api_rust",
|
"etke_openai_api_rust",
|
||||||
"matrix-sdk",
|
"matrix-sdk",
|
||||||
"mxidwc",
|
"mxidwc",
|
||||||
@@ -506,6 +507,15 @@ dependencies = [
|
|||||||
"zeroize",
|
"zeroize",
|
||||||
]
|
]
|
||||||
|
|
||||||
|
[[package]]
|
||||||
|
name = "chrono"
|
||||||
|
version = "0.4.38"
|
||||||
|
source = "registry+https://github.com/rust-lang/crates.io-index"
|
||||||
|
checksum = "a21f936df1771bf62b77f047b726c4625ff2e8aa607c01ec06e5a05bd8463401"
|
||||||
|
dependencies = [
|
||||||
|
"num-traits",
|
||||||
|
]
|
||||||
|
|
||||||
[[package]]
|
[[package]]
|
||||||
name = "cipher"
|
name = "cipher"
|
||||||
version = "0.4.4"
|
version = "0.4.4"
|
||||||
|
|||||||
@@ -19,6 +19,7 @@ anthropic-rs = "0.1.*"
|
|||||||
anyhow = "1.0.*"
|
anyhow = "1.0.*"
|
||||||
async-openai = "0.24.*"
|
async-openai = "0.24.*"
|
||||||
base64 = "0.22.*"
|
base64 = "0.22.*"
|
||||||
|
chrono = { version = "0.4.*", default-features = false, features = ["std", "now"] }
|
||||||
# We'd rather not depend on this, but we cannot use the ruma-events EventContent macro without it.
|
# We'd rather not depend on this, but we cannot use the ruma-events EventContent macro without it.
|
||||||
matrix-sdk = { version = "0.7.1", default-features = false }
|
matrix-sdk = { version = "0.7.1", default-features = false }
|
||||||
mxidwc = "1.0.*"
|
mxidwc = "1.0.*"
|
||||||
|
|||||||
@@ -68,6 +68,19 @@ Where appropriate, you'll mention best practices and common pitfalls.
|
|||||||
|
|
||||||
A prompt override can also be set globally, see [🛠️ Room Settings](./README.md#room-settings).
|
A prompt override can also be set globally, see [🛠️ Room Settings](./README.md#room-settings).
|
||||||
|
|
||||||
|
Prompts may contain the following **placeholder variables** which will be replaced *every time* the bot is interacted with:
|
||||||
|
|
||||||
|
| Placeholder | Description | Example |
|
||||||
|
|---------------------------|-------------|---------|
|
||||||
|
| `{{ baibot_name }}` | Name of the bot as configured in the `user.name` field in the [Static configuration](./README.md#static-configuration) | `Baibot` |
|
||||||
|
| `{{ baibot_model_id }}` | Text-Generation model ID as configured in the [🤖 agent](../agents.md)'s configuration | `gpt-4o` |
|
||||||
|
| `{{ baibot_now_utc }}` | Current date and time in UTC | `2024-09-20 (Friday), 14:26:42 UTC (local timezone/time: unknown)` |
|
||||||
|
|
||||||
|
Here's a prompt that combines some of the above variables:
|
||||||
|
|
||||||
|
> You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
|
|
||||||
|
|
||||||
### 🌡️ Temperature Override
|
### 🌡️ Temperature Override
|
||||||
|
|
||||||
You can override the [temperature](https://blogs.novita.ai/what-are-large-language-model-settings-temperature-top-p-and-max-tokens/#what-is-llm-temperature) (randomness / creativity) parameter configured at the [🤖 agent](../agents.md) level.
|
You can override the [temperature](https://blogs.novita.ai/what-are-large-language-model-settings-temperature-top-p-and-max-tokens/#what-is-llm-temperature) (randomness / creativity) parameter configured at the [🤖 agent](../agents.md) level.
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: https://api.anthropic.com/v1
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: claude-3-5-sonnet-20240620
|
model_id: claude-3-5-sonnet-20240620
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 8192
|
max_response_tokens: 8192
|
||||||
max_context_tokens: 204800
|
max_context_tokens: 204800
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: https://api.groq.com/openai/v1
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: llama3-70b-8192
|
model_id: llama3-70b-8192
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 4096
|
max_response_tokens: 4096
|
||||||
max_context_tokens: 131072
|
max_context_tokens: 131072
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: http://my-localai-self-hosted-service:8080/v1
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: gpt-4
|
model_id: gpt-4
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 4096
|
max_response_tokens: 4096
|
||||||
max_context_tokens: 128000
|
max_context_tokens: 128000
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: https://api.mistral.ai/v1
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: mistral-large-latest
|
model_id: mistral-large-latest
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 4096
|
max_response_tokens: 4096
|
||||||
max_context_tokens: 128000
|
max_context_tokens: 128000
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: http://my-ollama-self-hosted-service:11434/v1
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: gemma2:2b
|
model_id: gemma2:2b
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 4096
|
max_response_tokens: 4096
|
||||||
max_context_tokens: 128000
|
max_context_tokens: 128000
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: ''
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: some-model
|
model_id: some-model
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 4096
|
max_response_tokens: 4096
|
||||||
max_context_tokens: 128000
|
max_context_tokens: 128000
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: https://api.openai.com/v1
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: gpt-4o-2024-08-06
|
model_id: gpt-4o-2024-08-06
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 16384
|
max_response_tokens: 16384
|
||||||
max_context_tokens: 128000
|
max_context_tokens: 128000
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: https://openrouter.ai/api/v1
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: mattshumer/reflection-70b:free
|
model_id: mattshumer/reflection-70b:free
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 2048
|
max_response_tokens: 2048
|
||||||
max_context_tokens: 8192
|
max_context_tokens: 8192
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ base_url: https://api.together.xyz/v1
|
|||||||
api_key: YOUR_API_KEY_HERE
|
api_key: YOUR_API_KEY_HERE
|
||||||
text_generation:
|
text_generation:
|
||||||
model_id: meta-llama/Meta-Llama-3.1-405B-Instruct-Turbo
|
model_id: meta-llama/Meta-Llama-3.1-405B-Instruct-Turbo
|
||||||
prompt: You are a brief, but helpful bot.
|
prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
temperature: 1.0
|
temperature: 1.0
|
||||||
max_response_tokens: 2048
|
max_response_tokens: 2048
|
||||||
max_context_tokens: 8192
|
max_context_tokens: 8192
|
||||||
|
|||||||
@@ -73,7 +73,7 @@ agents:
|
|||||||
# api_key: ""
|
# api_key: ""
|
||||||
# text_generation:
|
# text_generation:
|
||||||
# model_id: gpt-4o-2024-08-06
|
# model_id: gpt-4o-2024-08-06
|
||||||
# prompt: You are a brief, but helpful bot.
|
# prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
# temperature: 1.0
|
# temperature: 1.0
|
||||||
# max_response_tokens: 16384
|
# max_response_tokens: 16384
|
||||||
# max_context_tokens: 128000
|
# max_context_tokens: 128000
|
||||||
@@ -97,7 +97,7 @@ agents:
|
|||||||
# api_key: null
|
# api_key: null
|
||||||
# text_generation:
|
# text_generation:
|
||||||
# model_id: gpt-4
|
# model_id: gpt-4
|
||||||
# prompt: You are a brief, but helpful bot.
|
# prompt: "You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
# temperature: 1.0
|
# temperature: 1.0
|
||||||
# max_response_tokens: 16384
|
# max_response_tokens: 16384
|
||||||
# max_context_tokens: 128000
|
# max_context_tokens: 128000
|
||||||
|
|||||||
@@ -19,3 +19,7 @@ pub use instantiation::Result as AgentInstantiationResult;
|
|||||||
|
|
||||||
pub use provider::{AgentProvider, AgentProviderInfo, ControllerTrait};
|
pub use provider::{AgentProvider, AgentProviderInfo, ControllerTrait};
|
||||||
pub use purpose::AgentPurpose;
|
pub use purpose::AgentPurpose;
|
||||||
|
|
||||||
|
pub(super) fn default_prompt() -> &'static str {
|
||||||
|
"You are a brief, but helpful bot called {{ baibot_name }} powered by the {{ baibot_model_id }} model. The date/time now is: {{ baibot_now_utc }}."
|
||||||
|
}
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ use serde::{Deserialize, Serialize};
|
|||||||
|
|
||||||
use anthropic_rs::models::claude::ClaudeModel;
|
use anthropic_rs::models::claude::ClaudeModel;
|
||||||
|
|
||||||
use crate::agent::provider::ConfigTrait;
|
use crate::agent::{default_prompt, provider::ConfigTrait};
|
||||||
|
|
||||||
#[derive(Debug, Clone, Serialize, Deserialize)]
|
#[derive(Debug, Clone, Serialize, Deserialize)]
|
||||||
pub struct Config {
|
pub struct Config {
|
||||||
@@ -58,7 +58,7 @@ impl Default for TextGenerationConfig {
|
|||||||
fn default() -> Self {
|
fn default() -> Self {
|
||||||
Self {
|
Self {
|
||||||
model_id: default_text_model_id(),
|
model_id: default_text_model_id(),
|
||||||
prompt: Some("You are a brief, but helpful bot.".to_owned()),
|
prompt: Some(default_prompt().to_owned()),
|
||||||
temperature: super::super::default_temperature(),
|
temperature: super::super::default_temperature(),
|
||||||
max_response_tokens: 8192,
|
max_response_tokens: 8192,
|
||||||
max_context_tokens: 204_800,
|
max_context_tokens: 204_800,
|
||||||
|
|||||||
@@ -95,11 +95,12 @@ impl ControllerTrait for Controller {
|
|||||||
));
|
));
|
||||||
};
|
};
|
||||||
|
|
||||||
let prompt_text = params
|
let prompt_text = params.prompt_variables.format(
|
||||||
.prompt_override
|
params
|
||||||
.unwrap_or(self.text_generation_prompt().unwrap_or("".to_owned()))
|
.prompt_override
|
||||||
.trim()
|
.unwrap_or(self.text_generation_prompt().unwrap_or("".to_owned()))
|
||||||
.to_owned();
|
.trim(),
|
||||||
|
);
|
||||||
|
|
||||||
let prompt_message = if prompt_text.is_empty() {
|
let prompt_message = if prompt_text.is_empty() {
|
||||||
None
|
None
|
||||||
@@ -225,20 +226,25 @@ impl ControllerTrait for Controller {
|
|||||||
}
|
}
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_generation_prompt(&self) -> Option<String> {
|
fn text_generation_model_id(&self) -> Option<String> {
|
||||||
let Some(text_generation_config) = &self.config.text_generation else {
|
self.config
|
||||||
return None;
|
.text_generation
|
||||||
};
|
.as_ref()
|
||||||
|
.map(|config| config.model_id.to_owned())
|
||||||
|
}
|
||||||
|
|
||||||
text_generation_config.prompt.clone()
|
fn text_generation_prompt(&self) -> Option<String> {
|
||||||
|
self.config
|
||||||
|
.text_generation
|
||||||
|
.as_ref()
|
||||||
|
.and_then(|config| config.prompt.clone())
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_generation_temperature(&self) -> Option<f32> {
|
fn text_generation_temperature(&self) -> Option<f32> {
|
||||||
let Some(text_generation_config) = &self.config.text_generation else {
|
self.config
|
||||||
return None;
|
.text_generation
|
||||||
};
|
.as_ref()
|
||||||
|
.map(|config| config.temperature)
|
||||||
Some(text_generation_config.temperature)
|
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_to_speech_voice(&self) -> Option<String> {
|
fn text_to_speech_voice(&self) -> Option<String> {
|
||||||
|
|||||||
@@ -13,6 +13,8 @@ pub trait ControllerTrait {
|
|||||||
|
|
||||||
fn ping(&self) -> impl std::future::Future<Output = anyhow::Result<PingResult>> + Send;
|
fn ping(&self) -> impl std::future::Future<Output = anyhow::Result<PingResult>> + Send;
|
||||||
|
|
||||||
|
fn text_generation_model_id(&self) -> Option<String>;
|
||||||
|
|
||||||
fn text_generation_prompt(&self) -> Option<String>;
|
fn text_generation_prompt(&self) -> Option<String>;
|
||||||
|
|
||||||
fn text_generation_temperature(&self) -> Option<f32>;
|
fn text_generation_temperature(&self) -> Option<f32>;
|
||||||
@@ -63,6 +65,14 @@ impl ControllerTrait for ControllerType {
|
|||||||
}
|
}
|
||||||
}
|
}
|
||||||
|
|
||||||
|
fn text_generation_model_id(&self) -> Option<String> {
|
||||||
|
match &self {
|
||||||
|
ControllerType::OpenAI(controller) => controller.text_generation_model_id(),
|
||||||
|
ControllerType::OpenAICompat(controller) => controller.text_generation_model_id(),
|
||||||
|
ControllerType::Anthropic(controller) => controller.text_generation_model_id(),
|
||||||
|
}
|
||||||
|
}
|
||||||
|
|
||||||
fn text_generation_prompt(&self) -> Option<String> {
|
fn text_generation_prompt(&self) -> Option<String> {
|
||||||
match &self {
|
match &self {
|
||||||
ControllerType::OpenAI(controller) => controller.text_generation_prompt(),
|
ControllerType::OpenAI(controller) => controller.text_generation_prompt(),
|
||||||
|
|||||||
@@ -9,5 +9,7 @@ pub use agent_provider::{AgentProvider, AgentProviderInfo};
|
|||||||
pub use image_generation::{ImageGenerationParams, ImageGenerationResult};
|
pub use image_generation::{ImageGenerationParams, ImageGenerationResult};
|
||||||
pub use ping::PingResult;
|
pub use ping::PingResult;
|
||||||
pub use speech_to_text::{SpeechToTextParams, SpeechToTextResult};
|
pub use speech_to_text::{SpeechToTextParams, SpeechToTextResult};
|
||||||
pub use text_generation::{TextGenerationParams, TextGenerationResult};
|
pub use text_generation::{
|
||||||
|
TextGenerationParams, TextGenerationPromptVariables, TextGenerationResult,
|
||||||
|
};
|
||||||
pub use text_to_speech::{TextToSpeechParams, TextToSpeechResult};
|
pub use text_to_speech::{TextToSpeechParams, TextToSpeechResult};
|
||||||
|
|||||||
@@ -1,8 +1,13 @@
|
|||||||
|
mod prompt_variables;
|
||||||
|
|
||||||
|
pub use prompt_variables::TextGenerationPromptVariables;
|
||||||
|
|
||||||
#[derive(Default)]
|
#[derive(Default)]
|
||||||
pub struct TextGenerationParams {
|
pub struct TextGenerationParams {
|
||||||
pub context_management_enabled: bool,
|
pub context_management_enabled: bool,
|
||||||
pub prompt_override: Option<String>,
|
pub prompt_override: Option<String>,
|
||||||
pub temperature_override: Option<f32>,
|
pub temperature_override: Option<f32>,
|
||||||
|
pub prompt_variables: TextGenerationPromptVariables,
|
||||||
}
|
}
|
||||||
|
|
||||||
pub struct TextGenerationResult {
|
pub struct TextGenerationResult {
|
||||||
@@ -0,0 +1,75 @@
|
|||||||
|
use chrono::{DateTime, Utc};
|
||||||
|
use std::collections::HashMap;
|
||||||
|
|
||||||
|
pub struct TextGenerationPromptVariables {
|
||||||
|
map: HashMap<String, String>,
|
||||||
|
}
|
||||||
|
|
||||||
|
impl Default for TextGenerationPromptVariables {
|
||||||
|
fn default() -> Self {
|
||||||
|
Self::new("unnamed", "unknown-model", Utc::now())
|
||||||
|
}
|
||||||
|
}
|
||||||
|
|
||||||
|
impl TextGenerationPromptVariables {
|
||||||
|
pub fn new(bot_name: &str, model_id: &str, utc_time: DateTime<Utc>) -> Self {
|
||||||
|
let mut map = HashMap::new();
|
||||||
|
|
||||||
|
map.insert("baibot_name".to_string(), bot_name.to_string());
|
||||||
|
map.insert("baibot_model_id".to_string(), model_id.to_string());
|
||||||
|
map.insert("baibot_now_utc".to_string(), format_utc_time(utc_time));
|
||||||
|
|
||||||
|
Self { map }
|
||||||
|
}
|
||||||
|
|
||||||
|
pub fn format(&self, text: &str) -> String {
|
||||||
|
let mut formatted_text = text.to_string();
|
||||||
|
|
||||||
|
for (key, value) in &self.map {
|
||||||
|
let placeholder = format!("{{{{ {} }}}}", key);
|
||||||
|
formatted_text = formatted_text.replace(&placeholder, value);
|
||||||
|
}
|
||||||
|
|
||||||
|
formatted_text
|
||||||
|
}
|
||||||
|
}
|
||||||
|
|
||||||
|
fn format_utc_time(time: DateTime<Utc>) -> String {
|
||||||
|
time.format("%Y-%m-%d (%A), %H:%M:%S UTC").to_string()
|
||||||
|
}
|
||||||
|
|
||||||
|
#[cfg(test)]
|
||||||
|
mod tests {
|
||||||
|
use super::*;
|
||||||
|
use chrono::{TimeZone, Timelike};
|
||||||
|
|
||||||
|
#[test]
|
||||||
|
fn test_new() {
|
||||||
|
// Intentionally injecting some sub-seconds to ensure formatting would ignore them.
|
||||||
|
let now_utc = Utc
|
||||||
|
.with_ymd_and_hms(2024, 9, 20, 18, 34, 15)
|
||||||
|
.unwrap()
|
||||||
|
.with_nanosecond(250000000)
|
||||||
|
.unwrap();
|
||||||
|
|
||||||
|
let variables = TextGenerationPromptVariables::new("baibot", "gpt-4o", now_utc);
|
||||||
|
|
||||||
|
assert_eq!(
|
||||||
|
variables.map.get("baibot_name"),
|
||||||
|
Some(&"baibot".to_string())
|
||||||
|
);
|
||||||
|
assert_eq!(
|
||||||
|
variables.map.get("baibot_model_id"),
|
||||||
|
Some(&"gpt-4o".to_string())
|
||||||
|
);
|
||||||
|
assert_eq!(
|
||||||
|
variables.map.get("baibot_now_utc"),
|
||||||
|
Some(&format_utc_time(now_utc))
|
||||||
|
);
|
||||||
|
|
||||||
|
let prompt = "Hello, I'm {{ baibot_name }} using {{ baibot_model_id }}. The date/time now is {{ baibot_now_utc }}.";
|
||||||
|
let expected = "Hello, I'm baibot using gpt-4o. The date/time now is 2024-09-20 (Friday), 18:34:15 UTC.";
|
||||||
|
|
||||||
|
assert_eq!(variables.format(prompt), expected);
|
||||||
|
}
|
||||||
|
}
|
||||||
@@ -21,5 +21,5 @@ pub use config::ConfigTrait;
|
|||||||
|
|
||||||
pub use entity::{
|
pub use entity::{
|
||||||
AgentProvider, AgentProviderInfo, ImageGenerationParams, PingResult, SpeechToTextParams,
|
AgentProvider, AgentProviderInfo, ImageGenerationParams, PingResult, SpeechToTextParams,
|
||||||
SpeechToTextResult, TextGenerationParams, TextToSpeechParams,
|
SpeechToTextResult, TextGenerationParams, TextGenerationPromptVariables, TextToSpeechParams,
|
||||||
};
|
};
|
||||||
|
|||||||
@@ -1,6 +1,6 @@
|
|||||||
use serde::{Deserialize, Serialize};
|
use serde::{Deserialize, Serialize};
|
||||||
|
|
||||||
use crate::agent::provider::ConfigTrait;
|
use crate::agent::{default_prompt, provider::ConfigTrait};
|
||||||
|
|
||||||
#[derive(Debug, Clone, Serialize, Deserialize)]
|
#[derive(Debug, Clone, Serialize, Deserialize)]
|
||||||
pub struct Config {
|
pub struct Config {
|
||||||
@@ -66,7 +66,7 @@ impl Default for TextGenerationConfig {
|
|||||||
fn default() -> Self {
|
fn default() -> Self {
|
||||||
Self {
|
Self {
|
||||||
model_id: default_text_model_id(),
|
model_id: default_text_model_id(),
|
||||||
prompt: Some("You are a brief, but helpful bot.".to_owned()),
|
prompt: Some(default_prompt().to_owned()),
|
||||||
temperature: super::super::default_temperature(),
|
temperature: super::super::default_temperature(),
|
||||||
max_response_tokens: 16_384,
|
max_response_tokens: 16_384,
|
||||||
max_context_tokens: 128_000,
|
max_context_tokens: 128_000,
|
||||||
|
|||||||
@@ -86,11 +86,12 @@ impl ControllerTrait for Controller {
|
|||||||
));
|
));
|
||||||
};
|
};
|
||||||
|
|
||||||
let prompt_text = params
|
let prompt_text = params.prompt_variables.format(
|
||||||
.prompt_override
|
params
|
||||||
.unwrap_or(self.text_generation_prompt().unwrap_or("".to_owned()))
|
.prompt_override
|
||||||
.trim()
|
.unwrap_or(self.text_generation_prompt().unwrap_or("".to_owned()))
|
||||||
.to_owned();
|
.trim(),
|
||||||
|
);
|
||||||
|
|
||||||
let prompt_message = if prompt_text.is_empty() {
|
let prompt_message = if prompt_text.is_empty() {
|
||||||
None
|
None
|
||||||
@@ -391,20 +392,25 @@ impl ControllerTrait for Controller {
|
|||||||
}
|
}
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_generation_prompt(&self) -> Option<String> {
|
fn text_generation_model_id(&self) -> Option<String> {
|
||||||
let Some(text_generation_config) = &self.config.text_generation else {
|
self.config
|
||||||
return None;
|
.text_generation
|
||||||
};
|
.as_ref()
|
||||||
|
.map(|config| config.model_id.to_owned())
|
||||||
|
}
|
||||||
|
|
||||||
text_generation_config.prompt.clone()
|
fn text_generation_prompt(&self) -> Option<String> {
|
||||||
|
self.config
|
||||||
|
.text_generation
|
||||||
|
.as_ref()
|
||||||
|
.and_then(|config| config.prompt.clone())
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_generation_temperature(&self) -> Option<f32> {
|
fn text_generation_temperature(&self) -> Option<f32> {
|
||||||
let Some(text_generation_config) = &self.config.text_generation else {
|
self.config
|
||||||
return None;
|
.text_generation
|
||||||
};
|
.as_ref()
|
||||||
|
.map(|config| config.temperature)
|
||||||
Some(text_generation_config.temperature)
|
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_to_speech_voice(&self) -> Option<String> {
|
fn text_to_speech_voice(&self) -> Option<String> {
|
||||||
|
|||||||
@@ -1,5 +1,6 @@
|
|||||||
use serde::{Deserialize, Serialize};
|
use serde::{Deserialize, Serialize};
|
||||||
|
|
||||||
|
use crate::agent::default_prompt;
|
||||||
use crate::agent::provider::openai::{
|
use crate::agent::provider::openai::{
|
||||||
ImageGenerationConfig as OpenAIImageGenerationConfig,
|
ImageGenerationConfig as OpenAIImageGenerationConfig,
|
||||||
SpeechToTextConfig as OpenAISpeechToTextConfig,
|
SpeechToTextConfig as OpenAISpeechToTextConfig,
|
||||||
@@ -75,7 +76,7 @@ impl Default for TextGenerationConfig {
|
|||||||
fn default() -> Self {
|
fn default() -> Self {
|
||||||
Self {
|
Self {
|
||||||
model_id: default_text_model_id(),
|
model_id: default_text_model_id(),
|
||||||
prompt: Some("You are a brief, but helpful bot.".to_owned()),
|
prompt: Some(default_prompt().to_owned()),
|
||||||
temperature: super::super::default_temperature(),
|
temperature: super::super::default_temperature(),
|
||||||
max_response_tokens: 4096,
|
max_response_tokens: 4096,
|
||||||
max_context_tokens: 128_000,
|
max_context_tokens: 128_000,
|
||||||
|
|||||||
@@ -84,11 +84,12 @@ impl ControllerTrait for Controller {
|
|||||||
));
|
));
|
||||||
};
|
};
|
||||||
|
|
||||||
let prompt_text = params
|
let prompt_text = params.prompt_variables.format(
|
||||||
.prompt_override
|
params
|
||||||
.unwrap_or(self.text_generation_prompt().unwrap_or("".to_owned()))
|
.prompt_override
|
||||||
.trim()
|
.unwrap_or(self.text_generation_prompt().unwrap_or("".to_owned()))
|
||||||
.to_owned();
|
.trim(),
|
||||||
|
);
|
||||||
|
|
||||||
let prompt_message = if prompt_text.is_empty() {
|
let prompt_message = if prompt_text.is_empty() {
|
||||||
None
|
None
|
||||||
@@ -409,20 +410,25 @@ impl ControllerTrait for Controller {
|
|||||||
}
|
}
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_generation_prompt(&self) -> Option<String> {
|
fn text_generation_model_id(&self) -> Option<String> {
|
||||||
let Some(text_generation_config) = &self.config.text_generation else {
|
self.config
|
||||||
return None;
|
.text_generation
|
||||||
};
|
.as_ref()
|
||||||
|
.map(|config| config.model_id.to_owned())
|
||||||
|
}
|
||||||
|
|
||||||
text_generation_config.prompt.clone()
|
fn text_generation_prompt(&self) -> Option<String> {
|
||||||
|
self.config
|
||||||
|
.text_generation
|
||||||
|
.as_ref()
|
||||||
|
.and_then(|config| config.prompt.clone())
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_generation_temperature(&self) -> Option<f32> {
|
fn text_generation_temperature(&self) -> Option<f32> {
|
||||||
let Some(text_generation_config) = &self.config.text_generation else {
|
self.config
|
||||||
return None;
|
.text_generation
|
||||||
};
|
.as_ref()
|
||||||
|
.map(|config| config.temperature)
|
||||||
Some(text_generation_config.temperature)
|
|
||||||
}
|
}
|
||||||
|
|
||||||
fn text_to_speech_voice(&self) -> Option<String> {
|
fn text_to_speech_voice(&self) -> Option<String> {
|
||||||
|
|||||||
@@ -4,7 +4,9 @@ use mxlink::{MatrixLink, MessageResponseType};
|
|||||||
|
|
||||||
use tracing::Instrument;
|
use tracing::Instrument;
|
||||||
|
|
||||||
use crate::agent::provider::{SpeechToTextParams, TextGenerationParams};
|
use crate::agent::provider::{
|
||||||
|
SpeechToTextParams, TextGenerationParams, TextGenerationPromptVariables,
|
||||||
|
};
|
||||||
use crate::agent::AgentInstance;
|
use crate::agent::AgentInstance;
|
||||||
use crate::agent::AgentPurpose;
|
use crate::agent::AgentPurpose;
|
||||||
use crate::agent::ControllerTrait;
|
use crate::agent::ControllerTrait;
|
||||||
@@ -405,6 +407,16 @@ async fn handle_stage_text_generation(
|
|||||||
|
|
||||||
let start_time = std::time::Instant::now();
|
let start_time = std::time::Instant::now();
|
||||||
|
|
||||||
|
let controller = agent.controller();
|
||||||
|
|
||||||
|
let prompt_variables = TextGenerationPromptVariables::new(
|
||||||
|
bot.name(),
|
||||||
|
&controller
|
||||||
|
.text_generation_model_id()
|
||||||
|
.unwrap_or("unknown-model".to_owned()),
|
||||||
|
chrono::Utc::now(),
|
||||||
|
);
|
||||||
|
|
||||||
let params = TextGenerationParams {
|
let params = TextGenerationParams {
|
||||||
context_management_enabled: message_context
|
context_management_enabled: message_context
|
||||||
.room_config_context()
|
.room_config_context()
|
||||||
@@ -417,10 +429,11 @@ async fn handle_stage_text_generation(
|
|||||||
temperature_override: message_context
|
temperature_override: message_context
|
||||||
.room_config_context()
|
.room_config_context()
|
||||||
.text_generation_temperature_override(),
|
.text_generation_temperature_override(),
|
||||||
|
|
||||||
|
prompt_variables,
|
||||||
};
|
};
|
||||||
|
|
||||||
let result = agent
|
let result = controller
|
||||||
.controller()
|
|
||||||
.generate_text(conversation, params)
|
.generate_text(conversation, params)
|
||||||
.instrument(span)
|
.instrument(span)
|
||||||
.await;
|
.await;
|
||||||
|
|||||||
Reference in New Issue
Block a user