2023-08-30 02:15:03 -04:00
|
|
|
#frozen_string_literal: true
|
|
|
|
|
|
|
|
module DiscourseAi
|
|
|
|
module AiBot
|
|
|
|
module Personas
|
|
|
|
class Persona
|
2024-01-04 08:44:07 -05:00
|
|
|
class << self
|
FEATURE: Add vision support to AI personas (Claude 3) (#546)
This commit adds the ability to enable vision for AI personas, allowing them to understand images that are posted in the conversation.
For personas with vision enabled, any images the user has posted will be resized to be within the configured max_pixels limit, base64 encoded and included in the prompt sent to the AI provider.
The persona editor allows enabling/disabling vision and has a dropdown to select the max supported image size (low, medium, high). Vision is disabled by default.
This initial vision support has been tested and implemented with Anthropic's claude-3 models which accept images in a special format as part of the prompt.
Other integrations will need to be updated to support images.
Several specs were added to test the new functionality at the persona, prompt building and API layers.
- Gemini is omitted, pending API support for Gemini 1.5. Current Gemini bot is not performing well, adding images is unlikely to make it perform any better.
- Open AI is omitted, vision support on GPT-4 it limited in that the API has no tool support when images are enabled so we would need to full back to a different prompting technique, something that would add lots of complexity
---------
Co-authored-by: Martin Brennan <martin@discourse.org>
2024-03-26 23:30:11 -04:00
|
|
|
def vision_enabled
|
|
|
|
false
|
|
|
|
end
|
|
|
|
|
|
|
|
def vision_max_pixels
|
|
|
|
1_048_576
|
|
|
|
end
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def system_personas
|
|
|
|
@system_personas ||= {
|
|
|
|
Personas::General => -1,
|
|
|
|
Personas::SqlHelper => -2,
|
|
|
|
Personas::Artist => -3,
|
|
|
|
Personas::SettingsExplorer => -4,
|
|
|
|
Personas::Researcher => -5,
|
|
|
|
Personas::Creative => -6,
|
|
|
|
Personas::DallE3 => -7,
|
2024-02-18 22:52:12 -05:00
|
|
|
Personas::DiscourseHelper => -8,
|
2024-03-07 14:37:23 -05:00
|
|
|
Personas::GithubHelper => -9,
|
2024-01-04 08:44:07 -05:00
|
|
|
}
|
|
|
|
end
|
2023-08-30 02:15:03 -04:00
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def system_personas_by_id
|
|
|
|
@system_personas_by_id ||= system_personas.invert
|
|
|
|
end
|
2023-09-14 02:46:56 -04:00
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def all(user:)
|
|
|
|
# listing tools has to be dynamic cause site settings may change
|
|
|
|
AiPersona.all_personas.filter do |persona|
|
|
|
|
next false if !user.in_any_groups?(persona.allowed_group_ids)
|
|
|
|
|
|
|
|
if persona.system
|
|
|
|
instance = persona.new
|
|
|
|
(
|
|
|
|
instance.required_tools == [] ||
|
|
|
|
(instance.required_tools - all_available_tools).empty?
|
|
|
|
)
|
|
|
|
else
|
|
|
|
true
|
|
|
|
end
|
|
|
|
end
|
|
|
|
end
|
2023-08-30 02:15:03 -04:00
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def find_by(id: nil, name: nil, user:)
|
|
|
|
all(user: user).find { |persona| persona.id == id || persona.name == name }
|
|
|
|
end
|
2023-12-07 16:42:56 -05:00
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def name
|
|
|
|
I18n.t("discourse_ai.ai_bot.personas.#{to_s.demodulize.underscore}.name")
|
|
|
|
end
|
2023-09-14 02:46:56 -04:00
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def description
|
|
|
|
I18n.t("discourse_ai.ai_bot.personas.#{to_s.demodulize.underscore}.description")
|
2023-08-30 02:15:03 -04:00
|
|
|
end
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def all_available_tools
|
|
|
|
tools = [
|
|
|
|
Tools::ListCategories,
|
|
|
|
Tools::Time,
|
|
|
|
Tools::Search,
|
|
|
|
Tools::Summarize,
|
|
|
|
Tools::Read,
|
|
|
|
Tools::DbSchema,
|
|
|
|
Tools::SearchSettings,
|
|
|
|
Tools::Summarize,
|
|
|
|
Tools::SettingContext,
|
2024-02-15 00:37:59 -05:00
|
|
|
Tools::RandomPicker,
|
2024-02-18 22:52:12 -05:00
|
|
|
Tools::DiscourseMetaSearch,
|
2024-03-07 14:37:23 -05:00
|
|
|
Tools::GithubFileContent,
|
|
|
|
Tools::GithubPullRequestDiff,
|
2024-03-28 01:01:58 -04:00
|
|
|
Tools::WebBrowser,
|
2024-01-04 08:44:07 -05:00
|
|
|
]
|
|
|
|
|
2024-03-07 14:37:23 -05:00
|
|
|
tools << Tools::GithubSearchCode if SiteSetting.ai_bot_github_access_token.present?
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
tools << Tools::ListTags if SiteSetting.tagging_enabled
|
|
|
|
tools << Tools::Image if SiteSetting.ai_stability_api_key.present?
|
|
|
|
|
|
|
|
tools << Tools::DallE if SiteSetting.ai_openai_api_key.present?
|
|
|
|
if SiteSetting.ai_google_custom_search_api_key.present? &&
|
|
|
|
SiteSetting.ai_google_custom_search_cx.present?
|
|
|
|
tools << Tools::Google
|
FEATURE: UI to update ai personas on admin page (#290)
Introduces a UI to manage customizable personas (admin only feature)
Part of the change was some extensive internal refactoring:
- AIBot now has a persona set in the constructor, once set it never changes
- Command now takes in bot as a constructor param, so it has the correct persona and is not generating AIBot objects on the fly
- Added a .prettierignore file, due to the way ALE is configured in nvim it is a pre-req for prettier to work
- Adds a bunch of validations on the AIPersona model, system personas (artist/creative etc...) are all seeded. We now ensure
- name uniqueness, and only allow certain properties to be touched for system personas.
- (JS note) the client side design takes advantage of nested routes, the parent route for personas gets all the personas via this.store.findAll("ai-persona") then child routes simply reach into this model to find a particular persona.
- (JS note) data is sideloaded into the ai-persona model the meta property supplied from the controller, resultSetMeta
- This removes ai_bot_enabled_personas and ai_bot_enabled_chat_commands, both should be controlled from the UI on a per persona basis
- Fixes a long standing bug in token accounting ... we were doing to_json.length instead of to_json.to_s.length
- Amended it so {commands} are always inserted at the end unconditionally, no need to add it to the template of the system message as it just confuses things
- Adds a concept of required_commands to stock personas, these are commands that must be configured for this stock persona to show up.
- Refactored tests so we stop requiring inference_stubs, it was very confusing to need it, added to plugin.rb for now which at least is clearer
- Migrates the persona selector to gjs
---------
Co-authored-by: Joffrey JAFFEUX <j.jaffeux@gmail.com>
Co-authored-by: Martin Brennan <martin@discourse.org>
2023-11-21 00:56:43 -05:00
|
|
|
end
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
tools
|
2023-08-30 02:15:03 -04:00
|
|
|
end
|
|
|
|
end
|
|
|
|
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
def id
|
|
|
|
@ai_persona&.id || self.class.system_personas[self.class]
|
|
|
|
end
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def tools
|
|
|
|
[]
|
2023-08-30 02:15:03 -04:00
|
|
|
end
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def required_tools
|
|
|
|
[]
|
|
|
|
end
|
2023-08-30 02:15:03 -04:00
|
|
|
|
2024-02-02 15:09:34 -05:00
|
|
|
def temperature
|
|
|
|
nil
|
|
|
|
end
|
|
|
|
|
|
|
|
def top_p
|
|
|
|
nil
|
|
|
|
end
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def options
|
|
|
|
{}
|
|
|
|
end
|
2023-08-30 02:15:03 -04:00
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def available_tools
|
|
|
|
self.class.all_available_tools.filter { |tool| tools.include?(tool) }
|
2023-08-30 02:15:03 -04:00
|
|
|
end
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
def craft_prompt(context)
|
|
|
|
system_insts =
|
|
|
|
system_prompt.gsub(/\{(\w+)\}/) do |match|
|
|
|
|
found = context[match[1..-2].to_sym]
|
|
|
|
found.nil? ? match : found.to_s
|
|
|
|
end
|
2023-08-30 02:15:03 -04:00
|
|
|
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
prompt_insts = <<~TEXT.strip
|
|
|
|
#{system_insts}
|
|
|
|
#{available_tools.map(&:custom_system_message).compact_blank.join("\n")}
|
|
|
|
TEXT
|
|
|
|
|
|
|
|
fragments_guidance = rag_fragments_prompt(context[:conversation_context].to_a)&.strip
|
|
|
|
|
|
|
|
if fragments_guidance.present?
|
|
|
|
if system_insts.include?("{uploads}")
|
|
|
|
prompt_insts = prompt_insts.gsub("{uploads}", fragments_guidance)
|
|
|
|
else
|
|
|
|
prompt_insts << fragments_guidance
|
|
|
|
end
|
|
|
|
end
|
|
|
|
|
2024-01-12 12:36:44 -05:00
|
|
|
prompt =
|
|
|
|
DiscourseAi::Completions::Prompt.new(
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
prompt_insts,
|
2024-01-12 12:36:44 -05:00
|
|
|
messages: context[:conversation_context].to_a,
|
2024-03-01 15:53:21 -05:00
|
|
|
topic_id: context[:topic_id],
|
|
|
|
post_id: context[:post_id],
|
2024-01-12 12:36:44 -05:00
|
|
|
)
|
2023-08-30 02:15:03 -04:00
|
|
|
|
FEATURE: Add vision support to AI personas (Claude 3) (#546)
This commit adds the ability to enable vision for AI personas, allowing them to understand images that are posted in the conversation.
For personas with vision enabled, any images the user has posted will be resized to be within the configured max_pixels limit, base64 encoded and included in the prompt sent to the AI provider.
The persona editor allows enabling/disabling vision and has a dropdown to select the max supported image size (low, medium, high). Vision is disabled by default.
This initial vision support has been tested and implemented with Anthropic's claude-3 models which accept images in a special format as part of the prompt.
Other integrations will need to be updated to support images.
Several specs were added to test the new functionality at the persona, prompt building and API layers.
- Gemini is omitted, pending API support for Gemini 1.5. Current Gemini bot is not performing well, adding images is unlikely to make it perform any better.
- Open AI is omitted, vision support on GPT-4 it limited in that the API has no tool support when images are enabled so we would need to full back to a different prompting technique, something that would add lots of complexity
---------
Co-authored-by: Martin Brennan <martin@discourse.org>
2024-03-26 23:30:11 -04:00
|
|
|
prompt.max_pixels = self.class.vision_max_pixels if self.class.vision_enabled
|
2024-01-12 12:36:44 -05:00
|
|
|
prompt.tools = available_tools.map(&:signature) if available_tools
|
|
|
|
|
|
|
|
prompt
|
FEATURE: UI to update ai personas on admin page (#290)
Introduces a UI to manage customizable personas (admin only feature)
Part of the change was some extensive internal refactoring:
- AIBot now has a persona set in the constructor, once set it never changes
- Command now takes in bot as a constructor param, so it has the correct persona and is not generating AIBot objects on the fly
- Added a .prettierignore file, due to the way ALE is configured in nvim it is a pre-req for prettier to work
- Adds a bunch of validations on the AIPersona model, system personas (artist/creative etc...) are all seeded. We now ensure
- name uniqueness, and only allow certain properties to be touched for system personas.
- (JS note) the client side design takes advantage of nested routes, the parent route for personas gets all the personas via this.store.findAll("ai-persona") then child routes simply reach into this model to find a particular persona.
- (JS note) data is sideloaded into the ai-persona model the meta property supplied from the controller, resultSetMeta
- This removes ai_bot_enabled_personas and ai_bot_enabled_chat_commands, both should be controlled from the UI on a per persona basis
- Fixes a long standing bug in token accounting ... we were doing to_json.length instead of to_json.to_s.length
- Amended it so {commands} are always inserted at the end unconditionally, no need to add it to the template of the system message as it just confuses things
- Adds a concept of required_commands to stock personas, these are commands that must be configured for this stock persona to show up.
- Refactored tests so we stop requiring inference_stubs, it was very confusing to need it, added to plugin.rb for now which at least is clearer
- Migrates the persona selector to gjs
---------
Co-authored-by: Joffrey JAFFEUX <j.jaffeux@gmail.com>
Co-authored-by: Martin Brennan <martin@discourse.org>
2023-11-21 00:56:43 -05:00
|
|
|
end
|
|
|
|
|
2024-03-01 15:53:21 -05:00
|
|
|
def find_tools(partial)
|
|
|
|
return [] if !partial.include?("</invoke>")
|
|
|
|
|
2024-01-04 08:44:07 -05:00
|
|
|
parsed_function = Nokogiri::HTML5.fragment(partial)
|
2024-03-01 15:53:21 -05:00
|
|
|
parsed_function.css("invoke").map { |fragment| find_tool(fragment) }.compact
|
|
|
|
end
|
|
|
|
|
|
|
|
protected
|
|
|
|
|
|
|
|
def find_tool(parsed_function)
|
2024-01-04 08:44:07 -05:00
|
|
|
function_id = parsed_function.at("tool_id")&.text
|
|
|
|
function_name = parsed_function.at("tool_name")&.text
|
2024-03-08 16:46:40 -05:00
|
|
|
return nil if function_name.nil?
|
2024-01-04 08:44:07 -05:00
|
|
|
|
|
|
|
tool_klass = available_tools.find { |c| c.signature.dig(:name) == function_name }
|
2024-03-08 16:46:40 -05:00
|
|
|
return nil if tool_klass.nil?
|
2024-01-04 08:44:07 -05:00
|
|
|
|
2024-01-04 22:39:32 -05:00
|
|
|
arguments = {}
|
|
|
|
tool_klass.signature[:parameters].to_a.each do |param|
|
|
|
|
name = param[:name]
|
|
|
|
value = parsed_function.at(name)&.text
|
|
|
|
|
|
|
|
if param[:type] == "array" && value
|
|
|
|
value =
|
|
|
|
begin
|
|
|
|
JSON.parse(value)
|
|
|
|
rescue JSON::ParserError
|
|
|
|
nil
|
|
|
|
end
|
|
|
|
end
|
|
|
|
|
|
|
|
arguments[name.to_sym] = value if value
|
|
|
|
end
|
2024-01-04 08:44:07 -05:00
|
|
|
|
|
|
|
tool_klass.new(
|
|
|
|
arguments,
|
2024-03-07 14:37:23 -05:00
|
|
|
tool_call_id: function_id || function_name,
|
2024-01-04 08:44:07 -05:00
|
|
|
persona_options: options[tool_klass].to_h,
|
|
|
|
)
|
2023-08-30 02:15:03 -04:00
|
|
|
end
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
|
|
|
|
def rag_fragments_prompt(conversation_context)
|
|
|
|
upload_refs =
|
|
|
|
UploadReference.where(target_id: id, target_type: "AiPersona").pluck(:upload_id)
|
|
|
|
|
|
|
|
return nil if !SiteSetting.ai_embeddings_enabled?
|
|
|
|
return nil if conversation_context.blank? || upload_refs.blank?
|
|
|
|
|
|
|
|
latest_interactions =
|
|
|
|
conversation_context
|
|
|
|
.select { |ctx| %i[model user].include?(ctx[:type]) }
|
|
|
|
.map { |ctx| ctx[:content] }
|
|
|
|
.last(10)
|
|
|
|
.join("\n")
|
|
|
|
|
|
|
|
strategy = DiscourseAi::Embeddings::Strategies::Truncation.new
|
|
|
|
vector_rep =
|
|
|
|
DiscourseAi::Embeddings::VectorRepresentations::Base.current_representation(strategy)
|
|
|
|
reranker = DiscourseAi::Inference::HuggingFaceTextEmbeddings
|
|
|
|
|
|
|
|
interactions_vector = vector_rep.vector_from(latest_interactions)
|
|
|
|
|
|
|
|
candidate_fragment_ids =
|
|
|
|
vector_rep.asymmetric_rag_fragment_similarity_search(
|
|
|
|
interactions_vector,
|
|
|
|
persona_id: id,
|
|
|
|
limit: reranker.reranker_configured? ? 50 : 10,
|
|
|
|
offset: 0,
|
|
|
|
)
|
|
|
|
|
2024-04-04 10:02:16 -04:00
|
|
|
fragments =
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
RagDocumentFragment.where(upload_id: upload_refs, id: candidate_fragment_ids).pluck(
|
|
|
|
:fragment,
|
2024-04-04 10:02:16 -04:00
|
|
|
:metadata,
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
)
|
|
|
|
|
|
|
|
if reranker.reranker_configured?
|
2024-04-04 10:02:16 -04:00
|
|
|
guidance = fragments.map { |fragment, _metadata| fragment }
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
ranks =
|
|
|
|
DiscourseAi::Inference::HuggingFaceTextEmbeddings
|
|
|
|
.rerank(conversation_context.last[:content], guidance)
|
|
|
|
.to_a
|
|
|
|
.take(10)
|
|
|
|
.map { _1[:index] }
|
|
|
|
|
|
|
|
if ranks.empty?
|
2024-04-04 10:02:16 -04:00
|
|
|
fragments = fragments.take(10)
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
else
|
2024-04-04 10:02:16 -04:00
|
|
|
fragments = ranks.map { |idx| fragments[idx] }
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
end
|
|
|
|
end
|
|
|
|
|
|
|
|
<<~TEXT
|
|
|
|
<guidance>
|
2024-04-04 10:02:16 -04:00
|
|
|
The following texts will give you additional guidance for your response.
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
We included them because we believe they are relevant to this conversation topic.
|
|
|
|
|
|
|
|
Texts:
|
|
|
|
|
2024-04-04 10:02:16 -04:00
|
|
|
#{
|
|
|
|
fragments
|
|
|
|
.map do |fragment, metadata|
|
|
|
|
if metadata.present?
|
|
|
|
["# #{metadata}", fragment].join("\n")
|
|
|
|
else
|
|
|
|
fragment
|
|
|
|
end
|
|
|
|
end
|
|
|
|
.join("\n")
|
|
|
|
}
|
FEATURE: AI Bot RAG support. (#537)
This PR lets you associate uploads to an AI persona, which we'll split and generate embeddings from. When building the system prompt to get a bot reply, we'll do a similarity search followed by a re-ranking (if available). This will let us find the most relevant fragments from the body of knowledge you associated with the persona, resulting in better, more informed responses.
For now, we'll only allow plain-text files, but this will change in the future.
Commits:
* FEATURE: RAG embeddings for the AI Bot
This first commit introduces a UI where admins can upload text files, which we'll store, split into fragments,
and generate embeddings of. In a next commit, we'll use those to give the bot additional information during
conversations.
* Basic asymmetric similarity search to provide guidance in system prompt
* Fix tests and lint
* Apply reranker to fragments
* Uploads filter, css adjustments and file validations
* Add placeholder for rag fragments
* Update annotations
2024-04-01 12:43:34 -04:00
|
|
|
</guidance>
|
|
|
|
TEXT
|
|
|
|
end
|
2023-08-30 02:15:03 -04:00
|
|
|
end
|
|
|
|
end
|
|
|
|
end
|
|
|
|
end
|