mirror of
https://github.com/meta-llama/llama-stack.git
synced 2025-10-16 06:53:47 +00:00
dataset validation
This commit is contained in:
parent
aefa84e70a
commit
7c280e18fb
4 changed files with 63 additions and 69 deletions
|
@ -62,10 +62,15 @@ class PandasDataframeDataset(BaseDataset):
|
|||
if self.df is None:
|
||||
self.load()
|
||||
|
||||
print(self.dataset_def.dataset_schema)
|
||||
# get columns names
|
||||
# columns = self.df[self.dataset_def.dataset_schema.keys()]
|
||||
print(self.df.columns)
|
||||
assert self.df is not None, "Dataset loading failed. Please check logs."
|
||||
|
||||
self.df = self.df[self.dataset_def.dataset_schema.keys()]
|
||||
|
||||
# check all columns in dataset schema are present
|
||||
assert len(self.df.columns) == len(self.dataset_def.dataset_schema)
|
||||
|
||||
# check all types match
|
||||
print(self.df.dtypes)
|
||||
|
||||
def load(self) -> None:
|
||||
if self.df is not None:
|
||||
|
|
Loading…
Add table
Add a link
Reference in a new issue